{"id":1227729,"url":"https://alion.io/job/prophecy-technologies-data-engineer","title":"Data Engineer","company":{"id":3800747,"name":"ProPhecy Technologies","domain":"prophecytechs.com","url":"https://alion.io/company/prophecy-technologies","size_band":null,"is_staffing_agency":false,"employer_type":"direct","is_intermediary":false,"listed_via":null,"ats_vendor":null,"truth_index":null},"role":"Data Science","role_family":"Data Science","seniority":"senior","employment_type":null,"work_mode":"on_site","remote_scope":null,"remote_scope_basis":null,"remote_working_hours":null,"hiring_geo_confidence":"structured","locations":["Hyderabad, India"],"countries":["IN"],"hiring_countries":[],"hiring_countries_total":0,"salary":null,"salary_estimate":{"min_usd":20000,"max_usd":41000,"period":"year","method":"role_seniority_country_remote_cell","sample_n":51},"experience_years_min":6,"visa_sponsorship":false,"relocation_package":false,"has_equity":false,"technologies":[{"name":"Agile","optional":false},{"name":"Amazon S3","optional":false},{"name":"Apache Kafka","optional":false},{"name":"AWS","optional":false},{"name":"AWS Glue","optional":false},{"name":"Azure","optional":false},{"name":"Azure Data Factory","optional":false},{"name":"CI/CD","optional":false},{"name":"Databricks","optional":false},{"name":"Delta Lake","optional":false},{"name":"ETL/ELT","optional":false},{"name":"Machine Learning","optional":false},{"name":"Python","optional":false},{"name":"Scala","optional":false},{"name":"Scrum","optional":false},{"name":"Spark","optional":false},{"name":"SQL","optional":false}],"status":"live","first_seen_at":"2026-09-23T08:02:56Z","employer_posted_date":null,"last_verified_at":"2026-09-23T08:02:56Z","board_verified":false,"closed_at":null,"days_open":7,"trust":{"level":"not_scored","repost_count":null,"flags":[],"days_open":7},"description":"Data Engineer with Scala\n\nJob Summary :\n\nWe are seeking a highly skilled and motivated Data Engineer with 6+ years of experience in designing, developing, and maintaining scalable data platforms and pipelines.\n\nThe ideal candidate will possess strong expertise in Apache Spark, Scala, Python, SQL, and Cloud Technologies (Azure/AWS), with the ability to independently own and deliver end-to-end data engineering solutions.\n\nThe role requires working closely with business stakeholders, data scientists, and engineering teams to build reliable, high-performance, and scalable data ecosystems that support analytics, reporting, and advanced data-driven initiatives.\n\nKey Responsibilities :\n\n- Design, develop, and maintain large-scale batch and real-time data pipelines using Apache Spark, Scala, and Python.\n\n- Build scalable and reliable data processing frameworks for ingesting, transforming, and integrating data from multiple sources.\n\n- Develop and optimize complex SQL queries, stored procedures, and data models to support reporting and analytics requirements.\n\n- Design and implement cloud-based data solutions using Azure and/or AWS services.\n\n- Create and manage ETL/ELT workflows using cloud-native tools such as Azure Data Factory, Azure Data Lake, AWS Glue, and Amazon S3.\n\n- Collaborate with business stakeholders, data analysts, and data scientists to understand requirements and deliver data solutions.\n\n- Monitor, troubleshoot, and optimize data pipelines to ensure performance, reliability, and data quality.\n\n- Implement best practices for data governance, security, scalability, and operational excellence.\n\n- Participate in code reviews, architecture discussions, and technical design sessions.\n\n- Support CI/CD implementation and automation of data engineering workflows.\n\n- Work with distributed data processing systems and contribute to platform modernization initiatives.\n\n- Mentor junior team members and contribute to knowledge-sharing activities within the team.\n\nExperience Required :\n\n- 6+ years of hands-on experience in Data Engineering, Data Warehousing, and Big Data technologies.\n\n- Strong experience developing scalable data pipelines using Apache Spark, Scala, and Python in enterprise environments.\n\n- Proven experience working with cloud platforms such as Microsoft Azure and/or AWS, including data storage, processing, and integration services.\n\n- Advanced knowledge of SQL, including complex query development, performance tuning, data modeling, and query optimization.\n\n- Experience designing and implementing end-to-end ETL/ELT workflows for large-scale data processing and analytics.\n\n- Demonstrated ability to independently own and deliver data engineering solutions from requirements gathering through deployment and production support.\n\n- Strong troubleshooting and problem-solving skills with experience resolving complex data quality, performance, and scalability challenges.\n\n- Experience working within Agile/Scrum teams and collaborating effectively with cross-functional stakeholders to deliver business-critical data solutions.\n\nPreferred to Have Skills :\n\n- Experience with Azure Data Factory (ADF), Azure Synapse Analytics, or AWS Glue.\n\n- Hands-on experience with Apache Kafka or other real-time streaming platforms.\n\n- Knowledge of Delta Lake, Databricks, or Lakehouse architectures.\n\n- Experience with CI/CD pipelines and DevOps practices for Data Engineering.\n\n- Familiarity with Data Governance, Data Quality, and Metadata Management frameworks.\n\n- Exposure to Generative AI, Machine Learning data pipelines, or Analytics platforms is an added advantage.\nSkills\nData Engineering, Scala, Python, Apache Spark, SQL, AWS, ETL, Azure Data Factory, AWS Glue","description_format":"text","description_chars":3712,"description_truncated":false,"requirements":{"experience_years_min":6,"management_years_min":null,"team_size_min":null,"manages_managers":false,"education":null,"security_clearance":false,"languages":[]},"benefits":[],"hiring_locations":[],"hiring_excludes":[],"relocation_offered":false,"industries":[],"lifecycle":[{"event":"open","at":"2026-09-25T13:06:44Z"}],"liveness":{"score":86,"band":"hot","label":"Hiring now","p_open":1,"p_active":0.864,"p_room":1,"age_days":6,"expected_fill_days":23,"reasons":["seen:6","velocity","win:early"],"computed_at":"2026-09-30T05:45:00Z"},"pay":null,"html_url":"https://alion.io/job/prophecy-technologies-data-engineer","json_url":"https://alion.io/job/prophecy-technologies-data-engineer.json","meta":{"generated_at":"2026-10-01T00:10:05Z","cache_seconds":300,"methodology":"https://alion.io/methodology","terms":"https://alion.io/terms","contact":"https://alion.io/contact","api":"https://alion.io/developers","usage":{"tier":"crawler","counted_by":"address","units_charged":1,"used_today":159,"day_limit":5000,"remaining_today":4841,"minute_limit":60,"resets_at":"2026-10-02T00:00:00Z"}}}