{"id":1318569,"url":"https://alion.io/job/cmcglobal-seniorlead-data-engineer","title":"(Senior/Lead) Data Engineer","company":{"id":3827704,"name":"CMC Global","domain":"cmcglobal.com.vn","url":"https://alion.io/company/cmcglobal","size_band":"201-500","is_staffing_agency":false,"employer_type":"staffing","is_intermediary":false,"listed_via":null,"ats_vendor":"Career site","truth_index":{"grade":"C","score":65,"open_postings":100,"ghost_share":0.59,"stale_share":0,"repost_share":0,"time_to_fill_p50_days":null,"computed_at":"2026-10-02T05:45:00Z"}},"role":"Data Science","role_family":"Data Science","seniority":"lead","employment_type":null,"work_mode":"on_site","remote_scope":null,"remote_scope_basis":null,"remote_working_hours":null,"hiring_geo_confidence":"structured","locations":["Vietnam"],"countries":["VN"],"hiring_countries":[],"hiring_countries_total":0,"salary":null,"salary_estimate":{"min_usd":21000,"max_usd":47000,"period":"year","method":"global_role_cell_scaled_by_country","sample_n":432},"experience_years_min":null,"visa_sponsorship":false,"relocation_package":false,"has_equity":false,"technologies":[{"name":"Apache Kafka","optional":false},{"name":"AWS","optional":false},{"name":"Azure","optional":false},{"name":"CI/CD","optional":false},{"name":"Databricks","optional":false},{"name":"Delta Lake","optional":false},{"name":"GCP","optional":false},{"name":"Hadoop","optional":false},{"name":"HBase","optional":false},{"name":"Informatica","optional":false},{"name":"Java","optional":false},{"name":"Machine Learning","optional":false},{"name":"Master Data Management","optional":false},{"name":"PostgreSQL","optional":false},{"name":"Python","optional":false},{"name":"Rest API","optional":false},{"name":"Scala","optional":false},{"name":"Scrum","optional":false},{"name":"Spark","optional":false},{"name":"SQL","optional":false},{"name":"Talend","optional":false}],"status":"live","first_seen_at":"2026-06-01T03:58:39Z","employer_posted_date":"2026-09-26","last_verified_at":"2026-09-30T06:28:02Z","board_verified":false,"closed_at":null,"days_open":124,"trust":{"level":"ghost","repost_count":0,"flags":["stale","company_stale"],"days_open":123},"description":"JOB DESCRIPTION\nCreate and manage a single master record for each business entity, ensuring data consistency, accuracy, and reliability.\nImplement data governance processes, including data quality management, data profiling, data remediation, and automated data lineage.\nCreate and maintain multiple robust and high-performance data processing pipelines within Cloud, Private Data Centre, and Hybrid data ecosystems.\nAssemble large, complex data sets from a wide variety of data sources.\nCollaborate with Data Scientists, Machine Learning Engineers, Business Analysts, and Business users to derive actionable insights and reliable foresights into customer acquisition, operational efficiency, and other key business performance metrics.\nDevelop, deploy, and maintain multiple microservices, REST APIs, and reporting services.\nDesign and implement internal processes to automate manual workflows, optimize data delivery, and re-design infrastructure for greater scalability.\nEstablish expertise in designing, analyzing, and troubleshooting large-scale distributed systems.\nSupport and work with cross-functional teams in a dynamic environment.\nREQUIREMENTS\nCritital: Data Engineer with strong Hadoop / Spark / Talend experience\nExperience building and operating large-scale data lakes and data warehouses.\nExperience with Hadoop ecosystem and big data tools, including Spark and Kafka.\nExperience with Master Data Management (MDM) tools and platforms such as Informatica MDM, Talend Data Catalog, Semarchy xDM, IBM PIM & IKC, or Profisee.\nFamiliarity with MDM processes such as golden record creation, survivorship, reconciliation, enrichment, and quality.\nExperience in data governance, including data quality management, data profiling, data remediation, and automated data lineage.\nExperience with stream-processing systems including Spark-Streaming.\nExperience working with Cloud services using one or more Cloud providers such as Azure, GCP, or AWS.\nExperience with Delta Lake and Databricks.\nAdvanced working experience with relational SQL and NoSQL databases, including Hive, HBase, and Postgres.\nDeep understanding of SQL and the ability to optimize data queries.\nExperience with object-oriented/object function scripting languages: Python, Java, Scala, etc.\nA successful history of manipulating, processing, and extracting value from large, disconnected datasets.\nExperience applying modern development principles (Scrum, TDD, continuous integration, and code reviews).\nProven ability to support and work with cross-functional teams in a dynamic environment.\nBENEFITS\nAttractive compensation package: 14-month salary scheme plus annual bonus and additional allowances\nAnnual bonus package tailored based on performance and contribution\nYoung, open, and dynamic working environment that promotes innovation and creativity\nOngoing learning and development with regular professional training and opportunities to enhance both technical and soft skills\nExposure to cutting-edge technologies and diverse real-world enterprise projects\nVibrant company culture with regular team-building activities, sports tournaments, arts events, Family Day, and more\nFull compliance with Vietnamese labor laws, plus additional internal perks such as annual company trips, special holidays, and other corporate benefits\nHOW TO APPLY\nPlease send your application via email: \n*By submitting your application to , you acknowledge that you have read, understood, and agreed to CMC Global’s REGULATIONS ON THE PROTECTION OF CANDIDATES’ PERSONAL INFORMATION.\nThe post (Senior/Lead) Data Engineer appeared first on CMC Global.","description_format":"text","description_chars":3654,"description_truncated":false,"requirements":{"experience_years_min":null,"management_years_min":null,"team_size_min":null,"manages_managers":false,"education":null,"security_clearance":false,"languages":[]},"benefits":[],"hiring_locations":[],"hiring_excludes":[],"relocation_offered":false,"industries":[],"lifecycle":[{"event":"open","at":"2026-09-26T20:13:45Z"}],"liveness":{"score":8,"band":"cold","label":"Long shot","p_open":1,"p_active":0.299,"p_room":0.28,"age_days":123,"expected_fill_days":40,"reasons":["conf:47","stale_co","velocity","ghost","win:tail","crowd:brand"],"computed_at":"2026-10-02T05:45:00Z"},"pay":null,"html_url":"https://alion.io/job/cmcglobal-seniorlead-data-engineer","json_url":"https://alion.io/job/cmcglobal-seniorlead-data-engineer.json","meta":{"generated_at":"2026-10-03T04:22:53Z","cache_seconds":300,"methodology":"https://alion.io/methodology","terms":"https://alion.io/terms","contact":"https://alion.io/contact","api":"https://alion.io/developers","usage":{"tier":"crawler","counted_by":"address","units_charged":1,"used_today":4634,"day_limit":5000,"remaining_today":366,"minute_limit":60,"resets_at":"2026-10-04T00:00:00Z"}}}