{"id":1230399,"url":"https://alion.io/job/infometry-senior-data-engineer","title":"Senior Data Engineer","company":{"id":3801196,"name":"Infometry","domain":"infometry.net","url":"https://alion.io/company/infometry","size_band":null,"is_staffing_agency":false,"employer_type":"direct","is_intermediary":false,"listed_via":null,"ats_vendor":null,"truth_index":null},"role":"Data Science","role_family":"Data Science","seniority":"senior","employment_type":null,"work_mode":"on_site","remote_scope":null,"remote_scope_basis":null,"remote_working_hours":null,"hiring_geo_confidence":"structured","locations":["Bengaluru, India"],"countries":["IN"],"hiring_countries":[],"hiring_countries_total":0,"salary":null,"salary_estimate":{"min_usd":22000,"max_usd":44000,"period":"year","method":"role_seniority_country_remote_cell","sample_n":51},"experience_years_min":8,"visa_sponsorship":false,"relocation_package":false,"has_equity":false,"technologies":[{"name":"Airflow","optional":false},{"name":"AWS","optional":false},{"name":"Azure","optional":false},{"name":"CI/CD","optional":false},{"name":"Databricks","optional":false},{"name":"Delta Lake","optional":false},{"name":"ETL/ELT","optional":false},{"name":"GCP","optional":false},{"name":"Git","optional":false},{"name":"NumPy","optional":false},{"name":"pySpark","optional":false},{"name":"Python","optional":false},{"name":"SQL","optional":false},{"name":"Spark","optional":true}],"status":"live","first_seen_at":"2026-09-18T13:57:12Z","employer_posted_date":null,"last_verified_at":"2026-09-18T13:57:12Z","board_verified":false,"closed_at":null,"days_open":16,"trust":{"level":"not_scored","repost_count":null,"flags":[],"days_open":16},"description":"Job Summary : \n\nWe are seeking an experienced Senior Data Engineer with strong expertise in Databricks, Apache Airflow, SQL, and Python to design and optimize data pipelines, ETL workflows, and cloud-based data solutions. The ideal candidate will have deep experience in big data processing, cloud-based analytics, and automation to drive efficient and scalable data engineering solutions.\n\nKey Responsibilities : \n\n- Design, develop, and optimize data pipelines using Databricks and Apache Airflow.\n\n- Implement PySpark-based transformations and processing in Databricks for handling large-scale data.\n\n- Develop and maintain SQL-based data pipelines, ensuring performance tuning and optimization.\n\n- Create Python scripts for automation, data transformation, and API-based data ingestion.\n\n- Work with Airflow DAGs to schedule and orchestrate data workflows efficiently.\n\n- Optimize data lake and data warehouse performance for scalability and reliability.\n\n- Integrate data pipelines with cloud platforms (AWS, Azure, or GCP) and various data storage solutions.\n\n- Ensure adherence to data security, governance, and compliance standards.\n\nRequired Skills & Qualifications : \n\n- 8-9 years of experience in Data Engineering or related fields.\n\n- Strong expertise in Databricks (PySpark, Delta Lake, DBSQL).\n\n- Proficiency in Apache Airflow for scheduling and orchestrating workflows.\n\n- Advanced SQL skills for data extraction, transformation, and performance tuning.\n\n- Strong programming skills in Python (pandas, NumPy, PySpark, APIs).\n\n- Experience with big data technologies and distributed computing.\n\n- Hands-on experience with cloud platforms (AWS / Azure / GCP).\n\n- Expertise in data warehousing and data modeling concepts.\n\n- Understanding of CI/CD pipelines, version control (Git)\n\n- Experience with data governance, security, and compliance best practices.\n\n- Excellent troubleshooting, debugging, and performance optimization skills.\nSkills\nData Engineering, Azure Databricks, Apache Airflow, SQL, Python, Data Pipeline, ETL, Data Warehousing, Data Modeling","description_format":"text","description_chars":2071,"description_truncated":false,"requirements":{"experience_years_min":8,"management_years_min":null,"team_size_min":null,"manages_managers":false,"education":null,"security_clearance":false,"languages":[]},"benefits":[],"hiring_locations":[],"hiring_excludes":[],"relocation_offered":false,"industries":["Analytics & BI Consulting","Data Engineering & Migration Services"],"lifecycle":[{"event":"open","at":"2026-09-25T14:00:00Z"}],"visa":[],"liveness":{"score":53,"band":"ok","label":"Likely open","p_open":0.85,"p_active":0.697,"p_room":0.9,"age_days":15,"expected_fill_days":24,"reasons":["seen:15","win:mid"],"computed_at":"2026-10-04T05:45:00Z"},"pay":null,"html_url":"https://alion.io/job/infometry-senior-data-engineer","json_url":"https://alion.io/job/infometry-senior-data-engineer.json","meta":{"generated_at":"2026-10-05T01:09:03Z","cache_seconds":300,"methodology":"https://alion.io/methodology","terms":"https://alion.io/terms","contact":"https://alion.io/contact","api":"https://alion.io/developers","usage":{"tier":"crawler","counted_by":"address","units_charged":1,"used_today":1457,"day_limit":5000,"remaining_today":3543,"minute_limit":60,"resets_at":"2026-10-06T00:00:00Z"}}}