{"id":720934,"url":"https://alion.io/job/shyftlabs-data-engineer-2","title":"Data Engineer","company":{"id":685982,"name":"Shyftlabs","domain":"shyftlabs.io","url":"https://alion.io/company/shyftlabs","size_band":null,"is_staffing_agency":false,"employer_type":"direct","is_intermediary":false,"listed_via":null,"ats_vendor":"Lever","truth_index":{"grade":"B","score":75,"open_postings":8,"ghost_share":0,"stale_share":1,"repost_share":0,"time_to_fill_p50_days":39,"computed_at":"2026-10-01T05:45:00Z"}},"role":"Data Science","role_family":"Data Science","seniority":"middle","employment_type":"full_time","work_mode":"on_site","remote_scope":null,"remote_scope_basis":null,"remote_working_hours":null,"hiring_geo_confidence":"structured","locations":["Noida, India"],"countries":["IN"],"hiring_countries":[],"hiring_countries_total":0,"salary":null,"salary_estimate":{"min_usd":12500,"max_usd":30000,"period":"year","method":"role_seniority_country_remote_cell","sample_n":13},"experience_years_min":3,"visa_sponsorship":false,"relocation_package":false,"has_equity":false,"technologies":[{"name":"Amazon Redshift","optional":false},{"name":"Anomaly Detection","optional":false},{"name":"Apache Kafka","optional":false},{"name":"AWS","optional":false},{"name":"Azure","optional":false},{"name":"BigQuery","optional":false},{"name":"ClickHouse","optional":false},{"name":"Docker","optional":false},{"name":"Druid","optional":false},{"name":"ETL/ELT","optional":false},{"name":"GCP","optional":false},{"name":"Google BigQuery","optional":false},{"name":"Java","optional":false},{"name":"Kubernetes","optional":false},{"name":"PostgreSQL","optional":false},{"name":"Python","optional":false},{"name":"Snowflake","optional":false},{"name":"SQL","optional":false},{"name":"Airflow","optional":true},{"name":"Looker","optional":true},{"name":"Spark","optional":true}],"status":"live","first_seen_at":"2026-07-10T07:03:56Z","employer_posted_date":"2026-07-10","last_verified_at":"2026-10-01T19:09:03Z","board_verified":true,"closed_at":null,"days_open":83,"trust":{"level":"ok","repost_count":null,"flags":[],"days_open":83},"description":"Position Overview:\nWe are seeking a skilled Data Engineer to design and build scalable data platforms that power analytics, reporting, and business-critical insights. You will develop high-performance batch and streaming data pipelines, optimize ETL/ELT workflows, and manage large-scale analytical databases while ensuring data reliability and performance. This role requires strong expertise in SQL, distributed data processing, cloud technologies, and event-driven architectures, along with close collaboration with Product, Analytics, and Backend teams to deliver robust data solutions.\nAt ShyftLabs, we live and breathe data. Since 2020, we’ve been helping Fortune 500 companies unlock growth with cutting-edge digital solutions that transform industries and create measurable business impact. We’re growing fast and we’re looking for passionate problem-solvers who are ready to turn big ideas into real outcomes.\nJob Responsibilities:\nDesign and build scalable batch and streaming data pipelines.\nDevelop and optimize ETL/ELT workflows using distributed data processing frameworks.\nOwn and optimize ClickHouse clusters for large-scale analytical workloads.\nDesign efficient data models for reporting and dashboarding use cases.\nBuild and maintain data ingestion pipelines from MongoDB, PostgreSQL, Kafka, APIs, and other data sources.\nImprove performance of large SQL workloads and analytical queries.\nBuild reliable monitoring, health checks, and data anomaly detection systems.\nWork closely with Product, Analytics, and Backend teams to deliver reliable reporting and insights. \nBasic Qualifications:\n3+ years of experience in Data Engineering.\nStrong SQL skills with expertise in query optimization.\nExperience with ClickHouse or other OLAP databases (BigQuery, Redshift, Snowflake, Druid, Pinot, etc.).\nStrong knowledge of PostgreSQL.\nExperience building ETL/ELT pipelines.\nProficiency in Java or Python.\nExperience with Apache Kafka and event-driven architectures.\nStrong understanding of data modeling and partitioning strategies.\nExperience working on cloud platforms (AWS/GCP/Azure).\nExperience with Docker and Kubernetes.\nKnowledge of monitoring and observability tools. \nPreferred Qualifications:\nExperience with Looker or BI platforms.\nExperience with Apache Spark or Dataproc.\nExperience with Airflow or workflow orchestration tools.\nUnderstanding of advertising technology (DSP, RTB, Attribution, Campaign Reporting).\nExperience with large-scale analytical systems processing billions of records.","description_format":"text","description_chars":2515,"description_truncated":false,"requirements":{"experience_years_min":3,"management_years_min":null,"team_size_min":null,"manages_managers":false,"education":null,"security_clearance":false,"languages":[]},"benefits":[],"hiring_locations":[],"hiring_excludes":[],"relocation_offered":false,"industries":["Artificial Intelligence","AI Consulting & Integration"],"lifecycle":[{"event":"open","at":"2026-09-11T07:11:31Z"}],"liveness":{"score":13,"band":"cold","label":"Long shot","p_open":1,"p_active":0.455,"p_room":0.28,"age_days":82,"expected_fill_days":39,"reasons":["conf:10","velocity","win:tail","crowd:"],"computed_at":"2026-10-01T05:45:00Z"},"pay":null,"html_url":"https://alion.io/job/shyftlabs-data-engineer-2","json_url":"https://alion.io/job/shyftlabs-data-engineer-2.json","meta":{"generated_at":"2026-10-02T00:41:58Z","cache_seconds":300,"methodology":"https://alion.io/methodology","terms":"https://alion.io/terms","contact":"https://alion.io/contact","api":"https://alion.io/developers","usage":{"tier":"crawler","counted_by":"address","units_charged":1,"used_today":663,"day_limit":5000,"remaining_today":4337,"minute_limit":60,"resets_at":"2026-10-03T00:00:00Z"}}}