{"id":1993597,"url":"https://alion.io/job/integrant-senior-lead-data-engineer","title":"Senior Lead Data Engineer","company":{"id":5272,"name":"Integrant","domain":"integrant.com","url":"https://alion.io/company/integrant","size_band":"201-500","is_staffing_agency":false,"employer_type":"direct","is_intermediary":false,"listed_via":null,"ats_vendor":"Workable","truth_index":{"grade":"A","score":89,"open_postings":28,"ghost_share":0.107,"stale_share":0,"repost_share":0,"time_to_fill_p50_days":91,"computed_at":"2026-10-08T05:49:30Z"}},"role":"Data Science","role_family":"Data Science","seniority":"lead","employment_type":"full_time","work_mode":"on_site","remote_scope":null,"remote_scope_basis":null,"remote_working_hours":null,"hiring_geo_confidence":"structured","locations":["Cairo, Egypt"],"countries":["EG"],"hiring_countries":[],"hiring_countries_total":0,"salary":null,"salary_estimate":{"min_usd":80000,"max_usd":180000,"period":"year","method":"global_role_cell_scaled_by_country","sample_n":521},"experience_years_min":null,"visa_sponsorship":false,"relocation_package":false,"has_equity":false,"technologies":[{"name":"AI Agents","optional":false},{"name":"Airflow","optional":false},{"name":"Amazon Kinesis","optional":false},{"name":"Amazon Redshift","optional":false},{"name":"Amazon S3","optional":false},{"name":"Apache HTTP Server","optional":false},{"name":"Apache Hudi","optional":false},{"name":"Apache Iceberg","optional":false},{"name":"Apache Kafka","optional":false},{"name":"AutoGen","optional":false},{"name":"AWS","optional":false},{"name":"AWS Glue","optional":false},{"name":"AWS Step Functions","optional":false},{"name":"Azure","optional":false},{"name":"Azure Cosmos DB","optional":false},{"name":"Azure Data Factory","optional":false},{"name":"Azure DevOps","optional":false},{"name":"BigQuery","optional":false},{"name":"Cassandra","optional":false},{"name":"CI/CD","optional":false},{"name":"Claude Code","optional":false},{"name":"Collibra","optional":false},{"name":"CrewAI","optional":false},{"name":"Cursor","optional":false},{"name":"Databricks","optional":false},{"name":"dbt","optional":false},{"name":"Delta Lake","optional":false},{"name":"Dimensional Modeling","optional":false},{"name":"DynamoDB","optional":false},{"name":"ETL/ELT","optional":false},{"name":"GCP","optional":false},{"name":"GitHub Actions","optional":false},{"name":"GitLab CI","optional":false},{"name":"Google BigQuery","optional":false},{"name":"Great Expectations","optional":false},{"name":"HashiCorp Vault","optional":false},{"name":"IAM","optional":false},{"name":"Informatica","optional":false},{"name":"Java","optional":false},{"name":"LangGraph","optional":false},{"name":"LlamaIndex","optional":false},{"name":"LLM","optional":false},{"name":"Looker","optional":false},{"name":"LookML","optional":false},{"name":"MS SQL","optional":false},{"name":"n8n","optional":false},{"name":"OpenAI Codex","optional":false},{"name":"PoC Library","optional":false},{"name":"Power BI","optional":false},{"name":"Python","optional":false},{"name":"Scala","optional":false},{"name":"Snowflake","optional":false},{"name":"Spark","optional":false},{"name":"SQL","optional":false},{"name":"SSAS","optional":false},{"name":"SSIS","optional":false},{"name":"Tableau","optional":false},{"name":"Talend","optional":false},{"name":"Teradata","optional":false},{"name":"Trino","optional":false},{"name":"Unity","optional":false},{"name":"Amazon SageMaker","optional":true},{"name":"AWS Bedrock","optional":true},{"name":"AWS Lambda","optional":true},{"name":"Docker","optional":true},{"name":"Kubernetes","optional":true},{"name":"LangChain","optional":true},{"name":"MLFlow","optional":true},{"name":"Model Context Protocol","optional":true},{"name":"Vertex AI","optional":true}],"status":"live","first_seen_at":"2026-10-07T08:00:37Z","employer_posted_date":"2026-10-07","last_verified_at":"2026-10-08T21:45:24Z","board_verified":true,"closed_at":null,"days_open":1,"trust":{"level":"ok","repost_count":0,"flags":["company_stale"],"days_open":1},"description":"Integrant is looking for game changers to join our team as a \"Senior Lead Data Engineer\". This is a senior, hands-on leader who designs data and AI solutions end to end and builds them. You'll be responsible for designing solutions for the ingestion, storage, processing, transformation, enrichment, and presentation of data for analytical consumption, and for personally implementing the critical parts of those solutions.\nEngage clients at multiple levels to elicit requirements, assess their current analytical challenges, and turn them into technical proposals and solution designs.\nSelect and customize analytical architectures (Data Warehouses, Data Lakes, Data Lakehouses, Data Fabrics, Data Meshes, etc.) and technically justify how they meet each client's needs.\nDesign and build agentic and LLM-powered solutions on top of modern data platforms.\nLead hands-on implementation of data pipelines, warehouses, and lakes, from design through performance tuning and production support.\nBuild data strategies, and coach clients and internal teams in the latest architectures and methodologies (DataOps, MLOps, etc.), including mentoring engineers on the team.\nRequirements\nArchitecture & Leadership\n10-12 years of experience in data engineering, including hands-on delivery at a senior level.\nBachelor's degree in Computer Science, Computer Engineering or other quantitative field. Post Graduate degrees preferable.\nThe ability to interact with stakeholders at multiple levels within a given organization, understand current analytical challenges and gather requirements.\nExcellent written and spoken English, with the ability to present solutions to technical and business audiences.\nProven experience writing technical proposals and solution designs for clients.\nA deep understanding of the evolution of analytical architectures from reporting databases to data warehouses, data lakes, data lakehouses, data fabrics, data meshes and beyond, with the ability to technically justify how a proposed architecture meets a client's analytical needs.\nExperience mentoring and guiding data engineers.\nHands-on Data Engineering\nExpert-level SQL and strong programming skills in Python (or Scala/Java) for data processing.\nExtensive experience with at least one enterprise data platform (e.g. Microsoft SQL Server stack, Oracle, Teradata).\nHands-on experience implementing data warehouses both on-premises and in the cloud (e.g. SQL Server, Oracle, Azure Synapse/Fabric, Amazon Redshift, Google BigQuery).\nStrong dimensional modeling: Dimensions, Facts, Slowly Changing Dimensions, Outriggers, Role-Playing, Junk, Degenerate and Multi-valued Dimensions; Transactional, Periodic Snapshot and Accumulating Snapshot Facts.\nHands-on experience implementing common ETL/ELT patterns using tools such as SSIS, Informatica, Talend, Azure Data Factory, AWS Glue, or code-based frameworks (e.g. Spark).\nHands-on experience implementing data quality, testing and observability for data pipelines (e.g. dbt tests, Great Expectations, Monte Carlo, Soda).\nExperience building semantic layers and OLAP models (e.g. SSAS, Power BI semantic models, LookML, dbt semantic layer).\nHands-on experience with workflow orchestration (e.g. Apache Airflow, Azure Data Factory, AWS Step Functions, Databricks Workflows).\nExperience setting up data lakes and lakehouses on cloud object storage (e.g. ADLS Gen2, Amazon S3, Google Cloud Storage) using open table formats (e.g. Delta Lake, Apache Iceberg, Apache Hudi).\nExperience optimizing query and workload performance and setting up high availability configurations.\nExperience administering and managing analytical solutions both on-premises and in the cloud.\nExperience building streaming pipelines (e.g. Kafka, Kinesis, Event Hubs, Google Pub/Sub, Spark Structured Streaming).\nExperience with query federation and data virtualization (e.g. Trino/Starburst, Amazon Athena, PolyBase, BigQuery federated queries).\nKnowledge of NoSQL databases (e.g. MongoDB, Cosmos DB, DynamoDB, Cassandra).\nKnowledge of BI tools (e.g. Power BI, Tableau, Looker, Qlik).\nCloud, Platforms & Governance\nExperience with multiple cloud data platforms (e.g. Microsoft Azure, AWS, Google Cloud).\nHands-on experience with Snowflake or Databricks.\nExperience estimating, monitoring and optimizing cloud data platform costs (e.g. warehouse sizing, compute/storage trade-offs, Snowflake credits, Databricks DBUs), and reflecting them in client proposals.\nExperience implementing data governance and security: cataloging, lineage, access control and secrets management (e.g. Unity Catalog, Microsoft Purview, AWS Lake Formation, Collibra, cloud IAM, HashiCorp Vault).\nExperience with CI/CD and version control for data solutions (e.g. GitHub Actions, GitLab CI, Azure DevOps).\nAI & Agentic\nHands-on exposure to LLM and agentic solutions, in PoC or production (e.g. LangGraph/AutoGen/CrewAI/LlamaIndex/n8n).\nExperience with AI-assisted development (e.g. Cursor/Codex/Claude Code).\nFamiliarity with DataOps.\nNice to Have\nFamiliarity with Agent Protocols (e.g. MCP).\nFamiliarity with transformation frameworks (e.g. dbt).\nFamiliarity with MLOps and ML platforms (e.g. MLflow, Azure ML, Amazon SageMaker, Vertex AI).\nFamiliarity with cloud AI services (e.g. Azure AI services, Amazon Bedrock, Vertex AI).\nFamiliarity with containerization tools and frameworks (Docker, Kubernetes, etc.).\nFamiliarity with the big data ecosystem, including Spark and Hive (e.g. on Databricks, Amazon EMR, Dataproc, HDInsight).\nFamiliarity with serverless functions for lightweight transformations (e.g. AWS Lambda, Azure Functions, Google Cloud Functions).\nBenefits\nSalary paid in USD\nCareer progression reviews every six months\nSupportive and friendly work environment\nPremium medical insurance [employee + family]\nSocial insurance\nEnglish language development courses\nInterest-free loans paid over 2.5 years\nTechnical development courses\nEmployment referral program\nPremium location in Maadi","description_format":"text","description_chars":5961,"description_truncated":false,"requirements":{"experience_years_min":null,"management_years_min":null,"team_size_min":null,"manages_managers":false,"education":{"level":"bachelor","optional":false},"security_clearance":false,"languages":[{"language":"English","level":"All levels","optional":false}]},"benefits":["Health insurance"],"hiring_locations":[],"hiring_excludes":[],"relocation_offered":false,"industries":["Artificial Intelligence","Mobile App Development Services","Web Development Services","AI Consulting & Integration"],"lifecycle":[{"event":"open","at":"2026-10-07T08:00:37Z"}],"visa":[],"liveness":{"score":86,"band":"hot","label":"Hiring now","p_open":1,"p_active":0.86,"p_room":1,"age_days":0,"expected_fill_days":91,"reasons":["conf:0","win:early"],"computed_at":"2026-10-08T05:49:30Z"},"pay":null,"html_url":"https://alion.io/job/integrant-senior-lead-data-engineer","json_url":"https://alion.io/job/integrant-senior-lead-data-engineer.json","meta":{"generated_at":"2026-10-09T00:24:32Z","cache_seconds":300,"methodology":"https://alion.io/methodology","terms":"https://alion.io/terms","contact":"https://alion.io/contact","api":"https://alion.io/developers","about":"Alion is a live layer of people, companies and AI agents: who they are, whether they are real and active right now, what they do and how to work with them, readable by people and by agents and paid per call.","catalog":"https://alion.io/catalog.json","usage":{"tier":"crawler","counted_by":"address","units_charged":1,"used_today":488,"day_limit":5000,"remaining_today":4512,"minute_limit":60,"resets_at":"2026-10-10T00:00:00Z"}},"offers":[{"id":"company.slices","title":"One company in depth, by slice","status":"live","price":{"credits":0.02,"usd":0.002,"plus_per_slice":{"credits":0.05,"usd":0.005}},"unit":"per company, plus each slice with data","note":"the employer in depth","call":{"mcp_tool":"get_company","arguments":{"id":5272},"rest":"https://alion.io/mcp/rest/get_company?id=5272"},"human":"https://alion.io/catalog?offer=company.slices&for=job%2Fintegrant-senior-lead-data-engineer"},{"id":"market.stats","title":"A market slice: pay, demand and time to fill","status":"live","price":{"credits":1,"usd":0.1},"unit":"per slice","note":"pay, demand and time to fill for this role and place","call":{"mcp_tool":"market_stats"},"human":"https://alion.io/catalog?offer=market.stats&for=job%2Fintegrant-senior-lead-data-engineer"},{"id":"job.search","title":"Open jobs by role, technology, place, pay and visa","status":"live","price":{"credits":0.02,"usd":0.002},"unit":"per posting in a list","note":"similar open postings","call":{"mcp_tool":"search_jobs"},"human":"https://alion.io/catalog?offer=job.search&for=job%2Fintegrant-senior-lead-data-engineer"},{"id":"company.verify","title":"Is this company real and active right now","status":"pilot","price":null,"unit":"per company","request":{"url":"https://alion.io/catalog/request","method":"POST","body":"{\"offer\": \"company.verify\", \"for\": \"job/integrant-senior-lead-data-engineer\", \"note\": \"what you need it for\"}"},"human":"https://alion.io/catalog?offer=company.verify&for=job%2Fintegrant-senior-lead-data-engineer"}]}