{"id":1316570,"url":"https://alion.io/job/onebeat-data-engineer-latam-3","title":"Data Engineer LATAM","company":{"id":2247241,"name":"Onebeat","domain":"onebeat.co","url":"https://alion.io/company/onebeat","size_band":"51-200","is_staffing_agency":false,"employer_type":"direct","is_intermediary":false,"listed_via":null,"ats_vendor":"Comeet","truth_index":null},"role":"Data Science","role_family":"Data Science","seniority":"middle","employment_type":"full_time","work_mode":"remote","remote_scope":"stated_countries","remote_scope_basis":"board_field","remote_working_hours":null,"hiring_geo_confidence":"structured","locations":["Mexico City, Mexico"],"countries":["MX"],"hiring_countries":["MX"],"hiring_countries_total":1,"salary":null,"salary_estimate":{"min_usd":34000,"max_usd":85000,"period":"year","method":"global_role_cell_scaled_by_country","sample_n":431},"experience_years_min":3,"visa_sponsorship":false,"relocation_package":false,"has_equity":false,"technologies":[{"name":"Amazon S3","optional":false},{"name":"Apache Iceberg","optional":false},{"name":"AWS","optional":false},{"name":"ClickHouse","optional":false},{"name":"ETL/ELT","optional":false},{"name":"Java","optional":false},{"name":"Python","optional":false},{"name":"Scala","optional":false},{"name":"Spark","optional":false},{"name":"SQL","optional":false},{"name":"CI/CD","optional":true},{"name":"Copilot","optional":true},{"name":"Cursor","optional":true},{"name":"Docker","optional":true},{"name":"Kubernetes","optional":true},{"name":"Looker","optional":true},{"name":"Metabase","optional":true},{"name":"Superset","optional":true},{"name":"Tableau","optional":true}],"status":"closed","first_seen_at":"2026-09-26T11:23:06Z","employer_posted_date":"2026-09-26","last_verified_at":"2026-09-27T11:22:17Z","board_verified":false,"closed_at":"2026-09-27T11:22:17Z","days_open":0,"trust":{"level":"not_scored","repost_count":null,"flags":[],"days_open":0},"description":"Description\nWe are seeking an experienced Data Engineer. The ideal candidate is self-motivated, can multitask, and is a proven team player. You will design, develop, manage, and maintain our open-source data platform, including our Data Lakehouse (S3, Apache Iceberg, and ClickHouse), ETL processes, and orchestration.\nWhat You Will Do\nDevelop a scalable data platform integrating multiple sources for easy access.\nDesign and enhance data tools (orchestration, governance, data lakehouse, BI, etc.).\nEnsure smooth operation of data systems for analysts, scientists, and engineers.\nOptimize data pipelines (ingestion, processing, and output) in a microservices environment.\nBuild, maintain, and monitor ETL/ELT processes .\nTroubleshoot and improve the performance, scalability, and reliability of the data infrastructure (S3, Apache Iceberg, ClickHouse).\nCollaborate cross-functionally with data scientists, analysts, and backend engineers to understand data needs and deliver solutions.\nImplement and champion data quality, governance, and security best practices across the platform.\nRequirements\n3+ years of experience as a Data Engineer or in a similar data infrastructure role.\nStrong proficiency in SQL and hands-on experience with data modeling.\nExperience with data lake/lakehouse architectures (e.g., Apache Iceberg, S3, or similar).\nExperience with analytical / columnar databases (e.g., ClickHouse or similar).\nExperience building and orchestrating ETL/ELT pipelines .\nStrong programming skills in Python and/or Scala/Java.\nExperience working within a microservices architecture and cloud environments (AWS preferred).\nSelf-motivated, strong multitasking skills, and a demonstrated team player.\nExcellent communication skills and the ability to work both independently and collaboratively.\nHands-on experience with Apache Spark (or similar technologies) for large-scale data processing.\nProfessional proficiency in written and spoken English.\nNote: this role is focused on batch data processing (not real-time streaming).\nNice to Have\nExperience working with and contributing to open-source data platforms and tools.\nFamiliarity with BI and visualization tools (e.g., Superset, Looker, Tableau, Metabase, or similar).\nExperience with containerization and orchestration (Docker, Kubernetes).\nExperience with infrastructure-as-code and CI/CD practices.\nExperience with AWS EMR and running Apache Spark workloads in a cloud environment.\nExperience leveraging AI-assisted development tools (e.g., GitHub Copilot, Cursor, or similar) to boost engineering productivity.","description_format":"text","description_chars":2573,"description_truncated":false,"requirements":{"experience_years_min":3,"management_years_min":null,"team_size_min":null,"manages_managers":false,"education":null,"security_clearance":false,"languages":[{"language":"English","level":"All levels","optional":false}]},"benefits":[],"hiring_locations":[{"name":"Mexico","iso":"MX","kind":"country"}],"hiring_excludes":[],"relocation_offered":false,"industries":["Artificial Intelligence","Cybersecurity","Travel Technology"],"lifecycle":[{"event":"open","at":"2026-09-26T18:59:58Z"},{"event":"close","at":"2026-09-27T11:22:17Z"}],"liveness":null,"pay":null,"html_url":"https://alion.io/job/onebeat-data-engineer-latam-3","json_url":"https://alion.io/job/onebeat-data-engineer-latam-3.json","meta":{"generated_at":"2026-09-29T03:50:27Z","cache_seconds":300,"methodology":"https://alion.io/methodology","terms":"https://alion.io/terms","contact":"https://alion.io/contact","api":"https://alion.io/developers","usage":{"tier":"crawler","counted_by":"address","units_charged":1,"used_today":3643,"day_limit":5000,"remaining_today":1357,"minute_limit":60,"resets_at":"2026-09-30T00:00:00Z"}}}