{"id":1272899,"url":"https://alion.io/job/tranzeal-data-engineer","title":"Data Engineer","company":{"id":3801539,"name":"Tranzeal","domain":"tranzeal.com","url":"https://alion.io/company/tranzeal","size_band":null,"is_staffing_agency":false,"employer_type":"direct","is_intermediary":false,"listed_via":null,"ats_vendor":null,"truth_index":null},"role":"Data Science","role_family":"Data Science","seniority":"senior","employment_type":null,"work_mode":"on_site","remote_scope":null,"remote_scope_basis":null,"remote_working_hours":null,"hiring_geo_confidence":"structured","locations":["Bengaluru, India"],"countries":["IN"],"hiring_countries":[],"hiring_countries_total":0,"salary":null,"salary_estimate":{"min_usd":22000,"max_usd":45000,"period":"year","method":"role_seniority_country_remote_cell","sample_n":51},"experience_years_min":7,"visa_sponsorship":false,"relocation_package":false,"has_equity":false,"technologies":[{"name":"Airflow","optional":false},{"name":"Amazon Kinesis","optional":false},{"name":"Amazon Redshift","optional":false},{"name":"Amazon S3","optional":false},{"name":"Apache Kafka","optional":false},{"name":"AWS","optional":false},{"name":"AWS Glue","optional":false},{"name":"AWS Step Functions","optional":false},{"name":"CI/CD","optional":false},{"name":"ETL/ELT","optional":false},{"name":"IAM","optional":false},{"name":"Machine Learning","optional":false},{"name":"pySpark","optional":false},{"name":"Spark","optional":false},{"name":"Python","optional":true}],"status":"live","first_seen_at":"2026-08-31T06:22:16Z","employer_posted_date":null,"last_verified_at":"2026-08-31T06:22:16Z","board_verified":false,"closed_at":null,"days_open":34,"trust":{"level":"not_scored","repost_count":null,"flags":[],"days_open":34},"description":"Job Description:\n\nWe're looking for a skilled Data Engineer to design, build, and maintain scalable data pipelines that power analytics, reporting, and machine learning initiatives. You'll work extensively with PySpark for large-scale data processing and the AWS ecosystem to build reliable, cloud-native data infrastructure.\n\nKey Responsibilities:\n\n- Design, develop, and maintain ETL/ELT pipelines using PySpark for batch and streaming data processing.\n\n- Build and manage data lake and data warehouse solutions on AWS (S3, Redshift, Glue, EMR, Athena, Lake Formation).\n\n- Develop and orchestrate workflows using AWS Step Functions, Apache Airflow, or AWS Glue Workflows.\n\n- Optimize Spark jobs for performance, cost, and scalability (partitioning, caching, cluster tuning).\n\n- Ingest data from multiple sources (APIs, databases, flat files, streaming platforms like Kafka/Kinesis).\n\n- Implement data quality checks, validation frameworks, and monitoring/alerting for pipeline health.\n\n- Collaborate with data analysts, data scientists, and business stakeholders to understand data requirements.\n\n- Design and maintain data models (star/snowflake schemas) for analytics use cases.\n\n- Write clean, well-documented, testable code following engineering best practices (CI/CD, version control).\n\n- Ensure data security, governance, and compliance (IAM policies, encryption, access controls).\n\n- Troubleshoot and resolve production data pipeline issues.\n\nSkills\nData Engineering, ETL, AWS, PySpark, Apache Airflow, Data Warehousing, DataLake, Data Validation","description_format":"text","description_chars":1555,"description_truncated":false,"requirements":{"experience_years_min":7,"management_years_min":null,"team_size_min":null,"manages_managers":false,"education":null,"security_clearance":false,"languages":[]},"benefits":[],"hiring_locations":[],"hiring_excludes":[],"relocation_offered":false,"industries":["IT Consulting & Digital Transformation","IT Staffing & Staff Augmentation"],"lifecycle":[{"event":"open","at":"2026-09-26T00:07:01Z"}],"visa":[],"liveness":{"score":19,"band":"cold","label":"Long shot","p_open":0.6,"p_active":0.588,"p_room":0.55,"age_days":33,"expected_fill_days":24,"reasons":["seen:33","win:tail"],"computed_at":"2026-10-04T05:45:00Z"},"pay":null,"html_url":"https://alion.io/job/tranzeal-data-engineer","json_url":"https://alion.io/job/tranzeal-data-engineer.json","meta":{"generated_at":"2026-10-04T16:37:54Z","cache_seconds":300,"methodology":"https://alion.io/methodology","terms":"https://alion.io/terms","contact":"https://alion.io/contact","api":"https://alion.io/developers","usage":{"tier":"search","counted_by":"address","units_charged":0,"used_today":0,"day_limit":null,"remaining_today":null,"minute_limit":null,"resets_at":"2026-10-05T00:00:00Z"}}}