{"id":1230702,"url":"https://alion.io/job/good-co-big-data-architect","title":"Big Data Architect","company":{"id":3801248,"name":"Good Co","domain":"goodco.co.in","url":"https://alion.io/company/good-co-2","size_band":null,"is_staffing_agency":true,"employer_type":"agency","is_intermediary":false,"listed_via":null,"ats_vendor":null,"truth_index":null},"role":"Data Science","role_family":"Data Science","seniority":"staff","employment_type":null,"work_mode":"on_site","remote_scope":null,"remote_scope_basis":null,"remote_working_hours":null,"hiring_geo_confidence":"structured","locations":["India"],"countries":["IN"],"hiring_countries":[],"hiring_countries_total":0,"salary":null,"salary_estimate":{"min_usd":39000,"max_usd":90000,"period":"year","method":"global_role_cell_scaled_by_country","sample_n":404},"experience_years_min":8,"visa_sponsorship":false,"relocation_package":false,"has_equity":false,"technologies":[{"name":"Airflow","optional":false},{"name":"Amazon Redshift","optional":false},{"name":"Apache Kafka","optional":false},{"name":"AWS","optional":false},{"name":"AWS Glue","optional":false},{"name":"Azure","optional":false},{"name":"Azure Data Factory","optional":false},{"name":"BigQuery","optional":false},{"name":"Databricks","optional":false},{"name":"ETL/ELT","optional":false},{"name":"GCP","optional":false},{"name":"Google BigQuery","optional":false},{"name":"Hadoop","optional":false},{"name":"Java","optional":false},{"name":"Python","optional":false},{"name":"Scala","optional":false},{"name":"Snowflake","optional":false},{"name":"Spark","optional":false},{"name":"SQL","optional":false}],"status":"live","first_seen_at":"2026-09-17T12:04:37Z","employer_posted_date":null,"last_verified_at":"2026-09-17T12:04:37Z","board_verified":false,"closed_at":null,"days_open":9,"trust":{"level":"not_scored","repost_count":null,"flags":[],"days_open":9},"description":"Role & Responsibilities :\n\n- Design and architect scalable, secure, and high-performance big data platforms capable of handling large volumes of structured and unstructured data.\n\n- Define end-to-end data architecture, data ingestion, processing, storage, integration, and analytics solutions.\n\n- Develop and implement data pipelines using technologies such as Apache Spark, Kafka, Hadoop, Hive, and Airflow.\n\n- Design cloud-based data platforms using AWS, Microsoft Azure, or Google Cloud Platform (GCP).\n\n- Architect modern data solutions using platforms such as Databricks, Snowflake, BigQuery, Amazon Redshift, or Azure Synapse.\n\n- Establish standards and best practices for data modeling, data governance, data quality, security, and metadata management.\n\n- Work closely with data engineers, data scientists, analysts, application teams, and business stakeholders to translate requirements into scalable technical solutions.\n\n- Evaluate existing data architecture and identify opportunities for modernization, optimization, automation, and cost reduction.\n\n- Design solutions for batch and real-time/streaming data processing.\n\n- Ensure high availability, fault tolerance, disaster recovery, and performance of data platforms.\n\n- Provide technical leadership and mentoring to data engineers and development teams.\n\nPreferred Candidate Profile :\n\n- 8 - 10 years of experience in data engineering, big data, cloud data platforms, or data architecture, with significant experience designing enterprise-scale data solutions.\n\n- Strong understanding of Big Data architecture and distributed computing concepts.\n\n- Hands-on experience with Apache Spark and at least one major big data ecosystem such as Hadoop/Hive.\n\n- Strong experience with Python, Java, or Scala.\n\n- Experience designing and implementing ETL/ELT and data pipelines.\n\n- Strong knowledge of Kafka or other event-streaming technologies for real-time data processing.\n\n- Experience with one or more cloud platforms: AWS, Azure, or GCP.\n\n- Experience with modern data platforms such as Databricks, Snowflake, BigQuery, Redshift, or Azure Synapse.\n\n- Strong knowledge of SQL, data modeling, data warehousing, data lakes, and lakehouse architecture.\n\n- Understanding of data governance, data quality, security, metadata, and master data concepts.\n\n- Experience with orchestration tools such as Airflow, Azure Data Factory, AWS Glue, or similar technologies.\n\n- Experience providing technical leadership, architecture guidance, and mentoring to engineering teams.\n\n- A Bachelor's or Master's degree in Computer Science, Information Technology, Engineering, or a related field is preferred.\n\nSkills\nBig Data, Big Data Architect, Data Engineering, Cloud, Apache Spark, Python, Java, ETL, Data Pipeline, SQL, Hadoop","description_format":"text","description_chars":2772,"description_truncated":false,"requirements":{"experience_years_min":8,"management_years_min":null,"team_size_min":null,"manages_managers":false,"education":null,"security_clearance":false,"languages":[]},"benefits":[],"hiring_locations":[],"hiring_excludes":[],"relocation_offered":false,"industries":[],"lifecycle":[{"event":"open","at":"2026-09-25T14:00:00Z"}],"liveness":{"score":73,"band":"hot","label":"Hiring now","p_open":1,"p_active":0.727,"p_room":1,"age_days":8,"expected_fill_days":31,"reasons":["seen:8","win:early"],"computed_at":"2026-09-26T05:45:00Z"},"pay":null,"html_url":"https://alion.io/job/good-co-big-data-architect","json_url":"https://alion.io/job/good-co-big-data-architect.json","meta":{"generated_at":"2026-09-27T01:28:53Z","cache_seconds":300,"methodology":"https://alion.io/methodology","terms":"https://alion.io/terms","contact":"https://alion.io/contact","api":"https://alion.io/developers","usage":{"tier":"crawler","counted_by":"address","units_charged":1,"used_today":1320,"day_limit":5000,"remaining_today":3680,"minute_limit":60,"resets_at":"2026-09-28T00:00:00Z"}}}