{"id":1625424,"url":"https://alion.io/job/mpulse-data-integration-engineer-ii","title":"Data Integration Engineer II","company":{"id":2083092,"name":"mPulse","domain":"mpulse.com","url":"https://alion.io/company/mpulse-com","size_band":"5000+","is_staffing_agency":false,"employer_type":"direct","is_intermediary":false,"listed_via":null,"ats_vendor":"Dayforce","truth_index":null},"role":"Data Science","role_family":"Data Science","seniority":"middle","employment_type":null,"work_mode":"remote","remote_scope":"stated_countries","remote_scope_basis":"board_field","remote_working_hours":null,"hiring_geo_confidence":"structured","locations":[],"countries":[],"hiring_countries":["US"],"hiring_countries_total":1,"salary":null,"salary_estimate":{"min_usd":93000,"max_usd":174000,"period":"year","method":"role_seniority_country_remote_cell","sample_n":53},"experience_years_min":3,"visa_sponsorship":false,"relocation_package":false,"has_equity":false,"technologies":[{"name":"Airflow","optional":false},{"name":"Amazon Redshift","optional":false},{"name":"Amazon S3","optional":false},{"name":"AWS","optional":false},{"name":"Bitbucket","optional":false},{"name":"CI/CD","optional":false},{"name":"dbt","optional":false},{"name":"ETL/ELT","optional":false},{"name":"GitHub","optional":false},{"name":"GitHub Actions","optional":false},{"name":"Jenkins","optional":false},{"name":"MS SQL","optional":false},{"name":"PostgreSQL","optional":false},{"name":"Python","optional":false},{"name":"Snowflake","optional":false},{"name":"SQL","optional":false},{"name":"Machine Learning","optional":true}],"status":"live","first_seen_at":"2026-09-02T05:00:00Z","employer_posted_date":"2026-10-01","last_verified_at":"2026-10-04T23:35:50Z","board_verified":true,"closed_at":null,"days_open":32,"trust":{"level":"ok","repost_count":null,"flags":[],"days_open":32},"description":"Job Summary\nWe’re looking for a passionate, self-motivated, and detail-oriented individual to join our Data Operations team. In this role, you will design, develop, and maintain scalable data pipelines that power analytics, reporting, and product capabilities within the predict vertical of the organization.\n\nYou will work closely with product engineering, implementation, analytics, and customer success teams to ensure data is accurate, reliable, and accessible. The ideal candidate has strong SQL and data warehousing experience, enjoys solving complex data problems, and is passionate about building reliable and well-documented data systems.\n\nDuties & Responsibilities\nDesign, develop, and maintain scalable data integration pipelines (ETL/ELT) to support data ingestion, transformation, cleansing, curation, and unification across multiple data sources.\nDevelop and support end-to-end data pipeline components, including ingestion, validation, transformation, cleansing, and curated data layer development.\nMonitor, maintain, and optimize production data pipelines to ensure reliability, performance, and successful execution of scheduled workflows.\nCreate and maintain comprehensive technical documentation for data pipelines, workflows, and data models.\nDevelop tools and frameworks to support automated data profiling, data quality monitoring, and unit testing to ensure high-quality and reliable data assets.\nCollaborate with implementation teams to identify, investigate, and resolve data anomalies during data onboarding and integration processes.\nPartner with product engineering teams to ensure accurate data capture and alignment with application data specifications and business requirements.\nProvide data-related support to analytics and customer success teams, assisting with troubleshooting, reporting needs, and client data inquiries.\n\nRequired Qualifications\nBachelor’s or master’s degree in computer science, Engineering, or a related technical field, or equivalent practical experience.\nMinimum of 3 years of professional experience in data engineering, data integration, or a related role.\nStrong proficiency in SQL, including complex querying, data transformation, and query performance optimization.\nExperience working with cloud-based data platforms and services, particularly within AWS (e.g., S3, Secrets Manager/Vault, DMS, or similar services).\nHands-on experience with modern data warehousing platforms, such as Snowflake, PostgreSQL, Amazon Redshift, or Microsoft SQL Server.\nExperience developing, debugging, and maintaining workflow orchestration pipelines using Apache Airflow, including DAG development and operational support.\nExperience using dbt (data build tool) to develop, test, and manage modular SQL-based data transformation models within modern data warehouse environments.\nExperience with version control systems and collaborative development workflows, using tools such as GitHub or Bitbucket.\nProficiency in Python, particularly for data manipulation, automation, and integration tasks.\nFamiliarity with CI/CD practices and automation tools, such as Jenkins or GitHub Actions.\nStrong written and verbal communication skills, with the ability to collaborate effectively across technical and non-technical teams.\n\nPreferred Qualifications\nExperience supporting data quality monitoring, data observability, or automated validation frameworks.\nFamiliarity with data science or machine learning workflows from a data engineering perspective.\nExperience working with healthcare-related datasets, such as claims, clinical, or regulatory data.\n\n*Please note, due to the requirements of this position, responses may automatically disqualify you from moving forward in the application process. Please review minimum qualifications thoroughly before applying. \n\nWhy Join our team?\nYou will have the opportunity to work on modern data platforms and tools, collaborate with cross-functional teams, and help build reliable data systems that drive meaningful insights and business value.\n\nWe’re as passionate about our people as we are about making our mark on healthcare. Fostering a fun and challenging environment that’s centered around personal and professional growth has brought us to where we are today. We are constantly seeking out new ways to reinvest in our team members because let’s face it, we all do our best work when we feel valued.","description_format":"text","description_chars":4385,"description_truncated":false,"requirements":{"experience_years_min":3,"management_years_min":null,"team_size_min":null,"manages_managers":false,"education":{"level":"bachelor","optional":true},"security_clearance":false,"languages":[]},"benefits":[],"hiring_locations":[{"name":"United States","iso":"US","kind":"country"}],"hiring_excludes":[],"relocation_offered":false,"industries":["Artificial Intelligence","Health Care","Professional Services"],"lifecycle":[{"event":"open","at":"2026-10-01T20:57:21Z"}],"visa":[],"liveness":{"score":23,"band":"cold","label":"Long shot","p_open":1,"p_active":0.665,"p_room":0.35,"age_days":32,"expected_fill_days":16,"reasons":["conf:0","velocity","win:tail"],"computed_at":"2026-10-04T05:45:00Z"},"pay":null,"html_url":"https://alion.io/job/mpulse-data-integration-engineer-ii","json_url":"https://alion.io/job/mpulse-data-integration-engineer-ii.json","meta":{"generated_at":"2026-10-05T00:28:43Z","cache_seconds":300,"methodology":"https://alion.io/methodology","terms":"https://alion.io/terms","contact":"https://alion.io/contact","api":"https://alion.io/developers","usage":{"tier":"crawler","counted_by":"address","units_charged":1,"used_today":534,"day_limit":5000,"remaining_today":4466,"minute_limit":60,"resets_at":"2026-10-06T00:00:00Z"}}}