{"id":1115850,"url":"https://alion.io/job/shyftlabs-data-engineer-databricks","title":"Data Engineer -Databricks","company":{"id":685982,"name":"Shyftlabs","domain":"shyftlabs.io","url":"https://alion.io/company/shyftlabs","size_band":null,"is_staffing_agency":false,"is_intermediary":false,"ats_vendor":"Lever","truth_index":{"grade":"B","score":75,"open_postings":8,"ghost_share":0,"stale_share":1,"repost_share":0,"time_to_fill_p50_days":39,"computed_at":"2026-09-23T05:45:00Z"}},"role":"Data Science","role_family":"Data Science","seniority":"middle","employment_type":"full_time","work_mode":"hybrid","remote_scope":null,"hiring_geo_confidence":"structured","locations":["Coimbatore, India"],"countries":["IN"],"hiring_countries":[],"hiring_countries_total":0,"salary":null,"salary_estimate":{"min_usd":20000,"max_usd":50000,"period":"year","method":"global_role_cell_scaled_by_country","sample_n":431},"experience_years_min":4,"visa_sponsorship":false,"relocation_package":false,"has_equity":false,"technologies":[{"name":"Amazon S3","optional":false},{"name":"AWS","optional":false},{"name":"CI/CD","optional":false},{"name":"Databricks","optional":false},{"name":"Delta Lake","optional":false},{"name":"Dimensional Modeling","optional":false},{"name":"ETL/ELT","optional":false},{"name":"Git","optional":false},{"name":"pySpark","optional":false},{"name":"Python","optional":false},{"name":"Rest API","optional":false},{"name":"Spark","optional":false},{"name":"SQL","optional":false},{"name":"Apache Kafka","optional":true},{"name":"dbt","optional":true}],"status":"live","first_seen_at":"2026-09-22T13:04:28Z","employer_posted_date":"2026-09-22","last_verified_at":"2026-09-24T00:12:19Z","board_verified":true,"closed_at":null,"days_open":1,"trust":{"level":"ok","repost_count":null,"flags":[],"days_open":1},"description":"Position Overview\nWe are looking for a Data Engineer with hands-on experience in building scalable data pipelines and data engineering solutions on the Databricks Lakehouse Platform. The ideal candidate should have strong expertise in Python, PySpark, SQL, Databricks, AWS, and REST API integrations for data ingestion, managing large volumes of data, and data export\nShyftLabs is a growing data product company that was founded in early 2020 and works primarily with Fortune 500 companies. We deliver digital solutions built to help accelerate the growth of businesses in various industries, by focusing on creating value through innovation.\nJob Responsibilities:\nDesign, develop, and maintain scalable ETL/ELT pipelines using Databricks,\nPySpark, and SQL.\nIntegrate data from multiple sources, including databases, Amazon S3, files, and REST APIs.\nBuild data pipelines with Databricks Unity Catalog.\nImplement business logic, data transformations, and dimensional data models.\nCreate, schedule, monitor, and optimize Databricks Jobs and Workflows.\nDesign and manage Delta Lake tables using Medallion Architecture (Bronze, Silver,Gold).\nEnsure data quality through validations, error handling, logging, and monitoring.\nOptimize Spark workloads for performance, scalability, and reliability.\nCollaborate with cross-functional teams to deliver production-ready data solutions.\nBasic Qualification:\nStrong expertise in Python, PySpark, and Advanced SQL.\nHands-on experience with the Databricks Lakehouse Platform.\nGood understanding of Unity Catalog, Delta Lake, Databricks Workflows/Jobs,\nClusters, Notebooks, Repos, and Medallion Architecture.\nExperience integrating with REST APIs for data ingestion and data export.\nStrong knowledge of ETL/ELT development, batch processing, incremental loading,\nand data transformation.\nExperience with data modeling (Star Schema, Snowflake Schema, Fact & Dimension\ntables, SCD concepts).\nUnderstanding of data warehousing concepts and best practices.\nExperience working with structured and semi-structured data (CSV, JSON, Parquet,\nDelta).\nKnowledge of partitioning, file optimization, Spark performance tuning, and query\noptimization.\nExperience with Git and CI/CD best practices\nPreferred Qualifications:\n4+ years of experience in Data Engineering with 2+ years of hands-on Databricks\nexperience.\nExperience with Auto Loader, Spark Declarative pipelines, Kafka, Airflow, or dbt is a plus.\nDatabricks certification is an added advantage.","description_format":"text","description_chars":2475,"description_truncated":false,"requirements":{"experience_years_min":4,"management_years_min":null,"team_size_min":null,"manages_managers":false,"education":null,"security_clearance":false,"languages":[]},"benefits":[],"hiring_locations":[{"name":"India","iso":"IN","kind":"country"}],"hiring_excludes":[],"relocation_offered":false,"industries":["Artificial Intelligence","AI Consulting & Integration"],"lifecycle":[{"event":"open","at":"2026-09-22T15:44:33Z"}],"liveness":{"score":90,"band":"hot","label":"Hiring now","p_open":1,"p_active":0.903,"p_room":1,"age_days":0,"expected_fill_days":39,"reasons":["conf:0","velocity","win:early"],"computed_at":"2026-09-23T05:45:00Z"},"pay":null,"html_url":"https://alion.io/job/shyftlabs-data-engineer-databricks","json_url":"https://alion.io/job/shyftlabs-data-engineer-databricks.json","meta":{"generated_at":"2026-09-24T02:00:11Z","cache_seconds":300,"methodology":"https://alion.io/methodology","terms":"https://alion.io/terms","contact":"https://alion.io/contact","api":"https://alion.io/developers"}}