{"id":1238852,"url":"https://alion.io/job/hiringblaze-data-engineer-ii","title":"Data Engineer II","company":{"id":3800132,"name":"HiringBlaze","domain":"hiringblaze.com","url":"https://alion.io/company/hiringblaze","size_band":null,"is_staffing_agency":true,"employer_type":"agency","is_intermediary":false,"listed_via":null,"ats_vendor":null,"truth_index":null},"role":"Data Science","role_family":"Data Science","seniority":"middle","employment_type":null,"work_mode":"hybrid","remote_scope":null,"remote_scope_basis":null,"remote_working_hours":null,"hiring_geo_confidence":"structured","locations":["Mumbai, India"],"countries":["IN"],"hiring_countries":[],"hiring_countries_total":0,"salary":null,"salary_estimate":{"min_usd":16000,"max_usd":39000,"period":"year","method":"role_seniority_country_remote_cell","sample_n":13},"experience_years_min":4,"visa_sponsorship":false,"relocation_package":false,"has_equity":false,"technologies":[{"name":"Apache Kafka","optional":false},{"name":"AWS","optional":false},{"name":"ETL/ELT","optional":false},{"name":"Hadoop","optional":false},{"name":"Machine Learning","optional":false},{"name":"Python","optional":false},{"name":"Scala","optional":false},{"name":"Spark","optional":false},{"name":"SQL","optional":false}],"status":"live","first_seen_at":"2026-09-25T14:21:34Z","employer_posted_date":null,"last_verified_at":"2026-09-25T14:21:34Z","board_verified":false,"closed_at":null,"days_open":3,"trust":{"level":"not_scored","repost_count":null,"flags":[],"days_open":3},"description":"Role Overview : \n\nAs a Data Engineer II based in Mumbai, you will serve as a cornerstone of our data infrastructure, architecting and maintaining robust systems that transform raw information into actionable business intelligence. \n\nYou will work closely with cross-functional teams, including data scientists, product managers, and software engineers, to design scalable pipelines that power our core analytical products. \n\nYour day-to-day will involve navigating complex data ecosystems to ensure high availability, performance, and data integrity, directly influencing how our leadership team makes data-driven decisions and how our customers experience our personalized service offerings.\n\nKey Responsibilities : \n\n- Design and implement end-to-end ETL/ELT pipelines to ingest, process, and store massive datasets, ensuring seamless data flow for downstream analytics and machine learning applications.\n\n- Optimize data warehouse performance and data modeling strategies to reduce query latency and improve cost-efficiency across our cloud infrastructure.\n\n- Collaborate with engineering stakeholders to integrate real-time data streaming solutions using Kafka, enabling low-latency insights for critical business operations.\n\n- Maintain and scale Big Data clusters, ensuring high availability and fault tolerance for mission-critical data processing tasks.\n\n- Mentor junior engineers through code reviews and technical guidance, fostering a culture of engineering excellence and best practices within the data team.\n\nRequired Skillset : \n\n- Demonstrated proficiency in Python and Scala for building complex data processing applications and automating infrastructure tasks.\n\n- Advanced expertise in SQL and NoSQL database management, with a proven ability to design efficient data models that support diverse analytical requirements.\n\n- Hands-on experience with Big Data frameworks including Spark and Hadoop, coupled with a strong understanding of distributed computing principles.\n\n- Proven track record of deploying and managing data solutions within AWS cloud environments, leveraging cloud-native tools to enhance scalability and security.\n\n- Strong communication skills with the ability to translate complex technical concepts into clear insights for non-technical stakeholders.\n\n- A collaborative mindset with the ability to thrive in a hybrid work environment, demonstrating adaptability and a proactive approach to problem-solving.\n\n- A Bachelors or Masters degree in Computer Science, Engineering, or a related quantitative field, supported by 4 - 7 years of professional experience in data engineering roles.\nSkills\nData Engineering, ETL, Machine Learning, Data Pipeline, Data Warehousing, Data Strategy, Data Analytics, Python, Scala, SQL, NoSQL","description_format":"text","description_chars":2761,"description_truncated":false,"requirements":{"experience_years_min":4,"management_years_min":null,"team_size_min":null,"manages_managers":false,"education":null,"security_clearance":false,"languages":[]},"benefits":["Hybrid work"],"hiring_locations":[],"hiring_excludes":[],"relocation_offered":false,"industries":[],"lifecycle":[{"event":"open","at":"2026-09-25T16:00:00Z"}],"liveness":{"score":44,"band":"fade","label":"Fading","p_open":1,"p_active":0.444,"p_room":1,"age_days":2,"expected_fill_days":8,"reasons":["seen:2","agency","velocity","win:early"],"computed_at":"2026-09-28T05:45:00Z"},"pay":null,"html_url":"https://alion.io/job/hiringblaze-data-engineer-ii","json_url":"https://alion.io/job/hiringblaze-data-engineer-ii.json","meta":{"generated_at":"2026-09-29T04:31:23Z","cache_seconds":300,"methodology":"https://alion.io/methodology","terms":"https://alion.io/terms","contact":"https://alion.io/contact","api":"https://alion.io/developers","usage":{"tier":"crawler","counted_by":"address","units_charged":1,"used_today":4285,"day_limit":5000,"remaining_today":715,"minute_limit":60,"resets_at":"2026-09-30T00:00:00Z"}}}