{"id":1254265,"url":"https://alion.io/job/combuilder-data-scientist","title":"Data Scientist","company":{"id":236844,"name":"Combuilder","domain":"combuilder.com.sg","url":"https://alion.io/company/combuilder","size_band":"1-10","is_staffing_agency":true,"employer_type":"agency","is_intermediary":false,"listed_via":null,"ats_vendor":null,"truth_index":null},"role":"Data Science","role_family":"Data Science","seniority":"junior","employment_type":"full_time","work_mode":"on_site","remote_scope":null,"remote_scope_basis":null,"remote_working_hours":null,"hiring_geo_confidence":"structured","locations":["Singapore"],"countries":["SG"],"hiring_countries":[],"hiring_countries_total":0,"salary":{"min":5500,"max":7500,"currency":"SGD","period":"month","gross":true,"usd_annual":70464},"salary_estimate":null,"experience_years_min":1,"visa_sponsorship":false,"relocation_package":false,"has_equity":false,"technologies":[{"name":"Agile","optional":false},{"name":"AI Agents","optional":false},{"name":"Amazon SageMaker","optional":false},{"name":"AWS","optional":false},{"name":"CI/CD","optional":false},{"name":"Git","optional":false},{"name":"LLM","optional":false},{"name":"NumPy","optional":false},{"name":"Pandas","optional":false},{"name":"pySpark","optional":false},{"name":"Python","optional":false},{"name":"RAG","optional":false},{"name":"Scikit-learn","optional":false},{"name":"SciPy","optional":false},{"name":"SQL","optional":false},{"name":"XGBoost","optional":false},{"name":"Spark","optional":true}],"status":"live","first_seen_at":"2026-09-16T00:00:00Z","employer_posted_date":null,"last_verified_at":"2026-09-16T00:00:00Z","board_verified":false,"closed_at":null,"days_open":19,"trust":{"level":"not_scored","repost_count":null,"flags":[],"days_open":19},"description":"The Data Scientist will design, develop and implement cohesive data integration and advanced analytics solutions involving both structured and unstructured data. The role supports mission‑critical initiatives, including predictive modelling, forecasting, operations research (optimization), text mining and network analytics, particularly within regulated and public‑sector environments.\n\nJob Responsibilities\nPrimarily responsible for applying the skills and knowledge gained about data analysis, analytics, data science to ensure the successful delivery of client engagements and initiatives within our Data and Analytics practice.\nDevelop and manage the end-to-end lifecycle of analytics projects from requirement gathering, data scoping, modelling to production (model deployment and monitoring).\nLead or support data requirement and analytics use-case workshops with business and technical stakeholders, translating business needs into clear analytical, data, and success metrics specifications.\nAttend and assist in facilitating project meetings / workshops with client stakeholders.\nPropose, implement, and validate data science models, ensuring functional and non-functional requirements such as explainability, fairness, scalability, security, integration and operational costs.\nParticipate actively in software development processes and best practices, documentation of requirements and software codes during the software development lifecycle.\nProduce high quality client-ready deliverables/document, with-ready-to-use content.\nIndependently drive assigned modules or workstreams with minimal supervision in a fast-paced project environment.\nPrepare user requirements, data development artefacts and technical documentation in accordance with governance and audit requirements.\nProactively research client business context, industry trends and functional domains, including public sector and government ecosystems, to stay current and relevant.\nContribute to the development of reusable project assets such as templates, analytical frameworks, processes, reports and presentation materials.\nPerform end-to-end testing and validation of migrated applications, including test design, execution, automation, defect management, and reporting to ensure successful migration outcomes.\nJob Requirements\nBachelor’s or master’s degree in Computer Science, Mathematics, Statistics, Business Analytics or equivalent.\n1-10 years’ experience in data science and data analytics fields.\nProven experience in data processing, feature selection, hyper-parameter optimization, model validation and visualization.\nProven experience in AWS SageMaker, Amazon Quick Sight, Python (e.g., Pandas, NumPy/SciPy, Scikit-Learn, XGBoost, pyspark, etc) and other related tools.\nExperience with Agentic AI/Generative AI (Large Language Model (LLM)–based solutions, including Retrieval-Augmented Generation, knowledge assistants, document intelligence, and conversational analytics), and modern data platform is an advantage.\nPreferred hands-on experience in data engineering, including data ingestion, transformation, pipeline development and working with data platforms or warehouses.\nStrong SQL skills with experience working on relational data models and large datasets.\nExperience in production software engineering routines such as test-driven development, code versioning with Git, conducting code reviews, and CI/CD.\nFamiliar with object-oriented programming concepts and their application to data science pipelines.\nDemonstrated ability to engage business stakeholders and lead data or analytics requirement workshops, translating complex business problems into actionable data solutions.\nDeep and eager interest in emerging technologies and the ability to leverage the technologies into solutions to meet our strategic and operational client needs.\nA self-starter with an analytical approach to problem solving.\nA client-centric, outcome driven and quality focused team player.\nDetailed oriented and is able to work in fast paced and agile environment.\nExcellent communication skills; both in written and spoken English.\nSkills\nObject-Oriented Programming, Retrieval-Augmented Generation (RAG), Liaising with project stakeholders, large datasets, SageMaker, Pipeline Development, Manage Software Development, Computer Science, Business Needs Analysis, Project Lifecycle Management, Requirements Documentation, Client Analysis, Work in A Fast Paced Environment, Functional Requirements, Document Delivery, Feature Selection","description_format":"text","description_chars":4511,"description_truncated":false,"requirements":{"experience_years_min":1,"management_years_min":null,"team_size_min":null,"manages_managers":false,"education":null,"security_clearance":false,"languages":[]},"benefits":[],"hiring_locations":[],"hiring_excludes":[],"relocation_offered":false,"industries":["Property Management"],"lifecycle":[{"event":"open","at":"2026-09-25T17:56:20Z"}],"visa":[],"liveness":{"score":3,"band":"cold","label":"Long shot","p_open":0.85,"p_active":0.086,"p_room":0.35,"age_days":18,"expected_fill_days":7,"reasons":["seen:18","agency","velocity","win:tail","comp:junior"],"computed_at":"2026-10-04T05:45:00Z"},"pay":{"stated_usd_annual":70464,"is_top_pay":false},"html_url":"https://alion.io/job/combuilder-data-scientist","json_url":"https://alion.io/job/combuilder-data-scientist.json","meta":{"generated_at":"2026-10-05T02:12:07Z","cache_seconds":300,"methodology":"https://alion.io/methodology","terms":"https://alion.io/terms","contact":"https://alion.io/contact","api":"https://alion.io/developers","usage":{"tier":"crawler","counted_by":"address","units_charged":1,"used_today":2899,"day_limit":5000,"remaining_today":2101,"minute_limit":60,"resets_at":"2026-10-06T00:00:00Z"}}}