{"id":2234608,"url":"https://alion.io/job/aldi-data-engineer-databricks-pyspark-sql-3","title":"Data Engineer (Databricks, PySpark, SQL)","company":{"id":2022379,"name":"ALDI Hungary","domain":"aldi.hu","url":"https://alion.io/company/aldi-hu","size_band":null,"is_staffing_agency":false,"employer_type":"direct","is_intermediary":false,"listed_via":null,"ats_vendor":"Career sitemap","truth_index":null},"role":"Data Science","role_family":"Data Science","seniority":"staff","employment_type":"full_time","work_mode":"remote","remote_scope":"stated_countries","remote_scope_basis":"board_field","remote_working_hours":null,"hiring_geo_confidence":"structured","locations":["Debrecen, Hungary"],"countries":["HU"],"hiring_countries":["HU"],"hiring_countries_total":1,"salary":null,"salary_estimate":{"min_usd":42000,"max_usd":102000,"period":"year","method":"global_role_cell_scaled_by_country","sample_n":631},"experience_years_min":10,"visa_sponsorship":false,"relocation_package":false,"has_equity":false,"technologies":[{"name":"Azure","optional":false},{"name":"Azure Data Factory","optional":false},{"name":"Databricks","optional":false},{"name":"Dimensional Modeling","optional":false},{"name":"ETL/ELT","optional":false},{"name":"Machine Learning","optional":false},{"name":"Python","optional":false},{"name":"Scala","optional":false},{"name":"Spark","optional":false},{"name":"SQL","optional":false},{"name":"Unix","optional":false},{"name":"pySpark","optional":true},{"name":"ServiceNow","optional":true}],"status":"live","first_seen_at":"2026-06-11T00:00:00Z","employer_posted_date":"2026-06-11","last_verified_at":"2026-10-11T19:43:58Z","board_verified":true,"closed_at":null,"days_open":122,"trust":{"level":"ghost","repost_count":0,"flags":["stale","company_stale"],"days_open":122},"description":"My responsibilities:\nDesign, develop, optimize, and maintain squad-specific data architectures and pipelines in line with ETL and data lake principles\nPrepare, align, and hand over data architecture and pipeline artefacts to the platform team for reuse across squads\nSolve complex technical data challenges that support business objectives\nDevelop data products for analytics, data scientists, and ML engineers to improve productivity within the team and across the organization\nAdvise, mentor, and support data and analytics professionals on data standards and best practices within the squad and organization\nContribute to the evaluation of emerging tools in data engineering and data science, helping to define and improve standards and ways of working\nDrive continuous improvement by contributing to training, capability development, and enhancements in analytical data engineering practices, standards, and processes\nThe knowledge I own:\nDegree in Computer Science or related field\n10+ years of experience in software or infrastructure development, including at least 4 years in data engineering within distributed computing, big data, or advanced analytics environments\nStrong expertise in SQL and data analysis, with proficiency in at least one programming language such as Python or Scala\nExperience in database development and data modeling, ideally with Databricks/Spark and SQL Server; knowledge of relational, NoSQL, and cloud-based databases.\nSolid understanding of distributed computing concepts, preferably with Spark or MapReduce\nHands-on experience with Azure tools such as Azure Data Factory, Azure Databricks, Event Hub, Synapse, and ML services\nGood understanding of data and analytics concepts, including dimensional modeling, ETL, data warehousing, reporting, data governance, and handling structured and unstructured data\nFamiliarity with Unix systems, especially shell scripting\nBasic knowledge of network concepts, connectivity, and troubleshooting\nFoundational understanding of machine learning, data science, AI, statistics, or applied mathematics\nStrong communication skills and fluent English; German is a plus\nTech Stack\nAzure Databricks\nAzure Data Factory\nPython\nPySpark\nServiceNow\nM365\nand other tools depending on the role\nThe offer that would convince me:\nWe develop and maintain our own product with high emphasis on quality and long-term stability\nCode quality matters: we follow Clean Code and SOLID principles backed by robust testing\nFlexible working hours and remote work options\nCompetitive salary with regular adjustments based on inflation and loyalty\nAccess to continuous learning via our internal learning platform and expert communities\nOnline application:\nPlease use our online application and attach your resume.\nAIIS Adatkezelési tájékoztató\nPrivacy notice","description_format":"text","description_chars":2805,"description_truncated":false,"requirements":{"experience_years_min":10,"management_years_min":null,"team_size_min":null,"manages_managers":false,"education":{"level":"bachelor","optional":false},"security_clearance":false,"languages":[{"language":"English","level":"Advanced (C1)","optional":false}]},"benefits":["Continuous learning","Flexible schedule"],"hiring_locations":[{"name":"Hungary","iso":"HU","kind":"country"}],"hiring_excludes":[],"relocation_offered":false,"industries":[],"lifecycle":[{"event":"open","at":"2026-10-10T23:25:30Z"}],"visa":[],"liveness":{"score":8,"band":"cold","label":"Long shot","p_open":1,"p_active":0.3,"p_room":0.28,"age_days":122,"expected_fill_days":27,"reasons":["conf:0","ghost","win:tail","crowd:brand"],"computed_at":"2026-10-11T20:42:48Z"},"pay":null,"html_url":"https://alion.io/job/aldi-data-engineer-databricks-pyspark-sql-3","json_url":"https://alion.io/job/aldi-data-engineer-databricks-pyspark-sql-3.json","meta":{"generated_at":"2026-10-11T20:42:48Z","cache_seconds":300,"methodology":"https://alion.io/methodology","terms":"https://alion.io/terms","contact":"https://alion.io/contact","api":"https://alion.io/developers","about":"Alion is a live layer of people, companies and AI agents: who they are, whether they are real and active right now, what they do and how to work with them, readable by people and by agents and paid per call.","catalog":"https://alion.io/catalog.json","usage":{"tier":"crawler_verified","counted_by":"address","units_charged":1,"used_today":9046,"day_limit":null,"remaining_today":null,"minute_limit":300,"resets_at":"2026-10-12T00:00:00Z"}},"offers":[{"id":"company.slices","title":"One company in depth, by slice","status":"live","price":{"credits":0.02,"usd":0.002,"plus_per_slice":{"credits":0.05,"usd":0.005}},"unit":"per company, plus each slice with data","note":"the employer in depth","call":{"mcp_tool":"get_company","arguments":{"id":2022379},"rest":"https://alion.io/mcp/rest/get_company?id=2022379"},"human":"https://alion.io/catalog?offer=company.slices&for=job%2Faldi-data-engineer-databricks-pyspark-sql-3"},{"id":"market.stats","title":"A market slice: pay, demand and time to fill","status":"live","price":{"credits":1,"usd":0.1},"unit":"per slice","note":"pay, demand and time to fill for this role and place","call":{"mcp_tool":"market_stats"},"human":"https://alion.io/catalog?offer=market.stats&for=job%2Faldi-data-engineer-databricks-pyspark-sql-3"},{"id":"job.search","title":"Open jobs by role, technology, place, pay and visa","status":"live","price":{"credits":0.02,"usd":0.002},"unit":"per posting in a list","note":"similar open postings","call":{"mcp_tool":"search_jobs"},"human":"https://alion.io/catalog?offer=job.search&for=job%2Faldi-data-engineer-databricks-pyspark-sql-3"},{"id":"company.verify","title":"Is this company real and active right now","status":"pilot","price":null,"unit":"per company","request":{"url":"https://alion.io/catalog/request","method":"POST","body":"{\"offer\": \"company.verify\", \"for\": \"job/aldi-data-engineer-databricks-pyspark-sql-3\", \"note\": \"what you need it for\"}"},"human":"https://alion.io/catalog?offer=company.verify&for=job%2Faldi-data-engineer-databricks-pyspark-sql-3"}]}