{"id":1238980,"url":"https://alion.io/job/teknuova-data-engineer","title":"Data Engineer","company":{"id":3802021,"name":"Teknuova","domain":"teknuova.com","url":"https://alion.io/company/teknuova","size_band":null,"is_staffing_agency":false,"employer_type":"direct","is_intermediary":false,"listed_via":null,"ats_vendor":null,"truth_index":null},"role":"Data Science","role_family":"Data Science","seniority":"senior","employment_type":null,"work_mode":"on_site","remote_scope":null,"remote_scope_basis":null,"remote_working_hours":null,"hiring_geo_confidence":"structured","locations":["Gurgaon, India","Bengaluru, India","Mumbai, India"],"countries":["IN"],"hiring_countries":[],"hiring_countries_total":0,"salary":null,"salary_estimate":{"min_usd":21000,"max_usd":43000,"period":"year","method":"role_seniority_country_remote_cell","sample_n":51},"experience_years_min":5,"visa_sponsorship":false,"relocation_package":false,"has_equity":false,"technologies":[{"name":"BigQuery","optional":false},{"name":"CI/CD","optional":false},{"name":"ETL/ELT","optional":false},{"name":"GCP","optional":false},{"name":"Git","optional":false},{"name":"Google BigQuery","optional":false},{"name":"Python","optional":false},{"name":"Spark","optional":false},{"name":"SQL","optional":false},{"name":"Airflow","optional":true},{"name":"Apache Kafka","optional":true},{"name":"dbt","optional":true},{"name":"Docker","optional":true},{"name":"IAM","optional":true},{"name":"Kubernetes","optional":true},{"name":"Terraform","optional":true}],"status":"live","first_seen_at":"2026-09-17T04:03:07Z","employer_posted_date":null,"last_verified_at":"2026-09-17T04:03:07Z","board_verified":false,"closed_at":null,"days_open":17,"trust":{"level":"not_scored","repost_count":null,"flags":[],"days_open":17},"description":"Job Description : \n\nWe are looking for a Data Engineer to design, build, and maintain scalable data pipelines and cloud-based data platforms on Google Cloud Platform. You will work with engineering, analytics, and business teams to ensure data is reliable, accessible, and optimized for downstream use cases.\n\nKey Responsibilities : \n\n- Design, develop, and maintain scalable data pipelines on Google Cloud Platform (GCP).\n\n- Build and manage batch and real-time data processing workflows.\n\n- Develop data ingestion, transformation, and integration solutions across multiple data sources.\n\n- Work with BigQuery to design data models, optimize queries, and improve data warehouse performance.\n\n- Build ETL/ELT pipelines using services such as Dataflow, Dataproc, Cloud Composer, and Pub/Sub.\n\n- Manage and process data stored in Google Cloud Storage.\n\n- Ensure data quality, consistency, security, and reliability across data platforms.\n\n- Monitor and troubleshoot data pipelines and production workflows.\n\n- Optimize cloud infrastructure and data processing workloads for performance and cost.\n\n- Collaborate with data scientists, analysts, software engineers, and business stakeholders.\n\n- Implement CI/CD and automation practices for data engineering workflows.\n\n- Maintain technical documentation for data pipelines, architecture, and processes.\n\nRequired Skills : \n\n- Strong hands-on experience with Google Cloud Platform.\n\n- Experience with BigQuery, Cloud Storage, Dataflow, Dataproc, Pub/Sub, and Cloud Composer.\n\n- Strong programming skills in Python.\n\n- Strong proficiency in SQL.\n\n- Experience designing and building ETL/ELT pipelines.\n\n- Good understanding of data warehousing, data modeling, and distributed data processing.\n\n- Experience working with large-scale structured and unstructured datasets.\n\n- Knowledge of Apache Beam, Spark, or similar data processing frameworks.\n\n- Familiarity with Git and version control.\n\n- Strong problem-solving and debugging skills.\n\nGood to Have : \n\n- Experience with Apache Airflow.\n\n- Knowledge of Terraform or Infrastructure as Code.\n\n- Experience with Docker and Kubernetes.\n\n- Familiarity with Kafka or other streaming technologies.\n\n- Exposure to dbt.\n\n- Experience implementing data governance and data security practices.\n\n- Knowledge of GCP IAM, networking, and cloud security.\n\n- Google Cloud certification such as Professional Data Engineer.\nSkills\nData Engineering, Google Cloud Platform, Data Modeling, BigQuery, Data Ingestion, Data Integration, ETL, Data Pipeline, Python, SQL, Data Warehousing","description_format":"text","description_chars":2559,"description_truncated":false,"requirements":{"experience_years_min":5,"management_years_min":null,"team_size_min":null,"manages_managers":false,"education":null,"security_clearance":false,"languages":[]},"benefits":[],"hiring_locations":[],"hiring_excludes":[],"relocation_offered":false,"industries":[],"lifecycle":[{"event":"open","at":"2026-09-25T16:00:00Z"}],"visa":[],"liveness":{"score":47,"band":"ok","label":"Likely open","p_open":0.85,"p_active":0.732,"p_room":0.75,"age_days":17,"expected_fill_days":24,"reasons":["seen:17","velocity","win:late"],"computed_at":"2026-10-04T05:45:00Z"},"pay":null,"html_url":"https://alion.io/job/teknuova-data-engineer","json_url":"https://alion.io/job/teknuova-data-engineer.json","meta":{"generated_at":"2026-10-05T02:57:49Z","cache_seconds":300,"methodology":"https://alion.io/methodology","terms":"https://alion.io/terms","contact":"https://alion.io/contact","api":"https://alion.io/developers","usage":{"tier":"crawler","counted_by":"address","units_charged":1,"used_today":4058,"day_limit":5000,"remaining_today":942,"minute_limit":60,"resets_at":"2026-10-06T00:00:00Z"}}}