{"id":1302534,"url":"https://alion.io/job/netapp-site-reliability-engineer-opensearch","title":"Site Reliability Engineer (OpenSearch)","company":{"id":58700,"name":"NetApp","domain":"netapp.com","url":"https://alion.io/company/netapp","size_band":"1001-5000","is_staffing_agency":false,"employer_type":"direct","is_intermediary":false,"listed_via":null,"ats_vendor":"Eightfold","truth_index":null},"role":"DevOps","role_family":"DevOps","seniority":"middle","employment_type":null,"work_mode":"on_site","remote_scope":null,"remote_scope_basis":null,"remote_working_hours":null,"hiring_geo_confidence":"structured","locations":["Bengaluru, India"],"countries":["IN"],"hiring_countries":[],"hiring_countries_total":0,"salary":null,"salary_estimate":{"min_usd":19000,"max_usd":57000,"period":"year","method":"global_role_cell_scaled_by_country","sample_n":618},"experience_years_min":4,"visa_sponsorship":false,"relocation_package":false,"has_equity":false,"technologies":[{"name":"AWS","optional":false},{"name":"Azure","optional":false},{"name":"Docker","optional":false},{"name":"GCP","optional":false},{"name":"Linux","optional":false},{"name":"OpenSearch","optional":false},{"name":"Apache Kafka","optional":true},{"name":"Bash","optional":true},{"name":"Git","optional":true},{"name":"Java","optional":true},{"name":"Jira","optional":true},{"name":"Python","optional":true}],"status":"live","first_seen_at":"2026-05-27T03:33:27Z","employer_posted_date":"2026-08-28","last_verified_at":"2026-09-26T22:31:07Z","board_verified":true,"closed_at":null,"days_open":122,"trust":{"level":"stale","repost_count":0,"flags":["stale"],"days_open":122},"description":"Job Summary\nNetApp is seeking a Technical Operations Engineer (OpenSearch) to join our growing Instaclustr team in Bangalore, India. In this role, you will be part of a frontline Site Reliability Engineering (SRE) team responsible for ensuring the availability, performance, and reliability of large-scale, cloud-hosted OpenSearch clusters.\nYou will work in a highly automated environment managing distributed open-source systems at scale, collaborating with global customers across industries such as banking, telecom, gaming, and technology. This role requires strong operational expertise, problem-solving skills, and a passion for learning and working with modern cloud-native and open-source technologies.\nJob Requirements\nProvide end-to-end operational support for OpenSearch clusters deployed across public cloud platforms (AWS, Azure, GCP).\nMonitor, troubleshoot, and resolve complex production issues, ensuring high availability and performance.\nPerform cluster lifecycle operations, including upgrades, migrations, maintenance, and scaling activities.\nParticipate in L2 on-call rotations, ensuring timely incident response and resolution.\nCollaborate with customer engineering teams to diagnose and resolve issues related to OpenSearch and other supported technologies.\nWork closely with internal teams to enhance reliability, automation, and operational efficiency.\nDevelop and improve automation tools, scripts, and operational processes.\nAnalyse system behaviour and proactively identify opportunities for performance optimisation and reliability improvements.\nContribute to knowledge sharing, documentation, and continuous improvement initiatives.\nRequired Skills & Experience\nHands-on experience with OpenSearch (including troubleshooting, upgrades, and migrations) or strong willingness to develop deep expertise.\nExperience with public cloud platforms such as AWS, Azure, or GCP.\nStrong Linux system administration skills and comfort with command-line environments.\nSolid understanding of distributed systems, networking, and OS internals.\nExperience with containerisation technologies (e.g., Docker).\nStrong problem-solving skills with the ability to debug complex production issues.\nExcellent communication skills (written and verbal) with a customer-focused mindset.\nAbility to work effectively in a collaborative, fast-paced environment and take ownership of tasks.\nPreferred Skills\nExperience working with other distributed systems such as Cassandra or Kafka.\nFamiliarity with source code debugging and issue investigation (e.g., Jira, codebase review).\nProgramming/scripting skills in Python, Java, or Bash.\nExperience with Git or version control systems.\nPrior experience in customer support or technical operations roles\nEducation\nTypically requires a minimum of 4-8 years of related experience with a Bachelor’s degree or 6 years and a Master’s degree; or a PhD with 3 years experience; or equivalent experience.","description_format":"text","description_chars":2938,"description_truncated":false,"requirements":{"experience_years_min":4,"management_years_min":null,"team_size_min":null,"manages_managers":false,"education":{"level":"bachelor","optional":false},"security_clearance":false,"languages":[]},"benefits":[],"hiring_locations":[],"hiring_excludes":[],"relocation_offered":false,"industries":["Cybersecurity","Natural Resources","Cloud Platforms (IaaS & PaaS)","Digital Storage"],"lifecycle":[{"event":"open","at":"2026-09-26T12:11:05Z"}],"liveness":{"score":8,"band":"cold","label":"Long shot","p_open":1,"p_active":0.284,"p_room":0.28,"age_days":122,"expected_fill_days":27,"reasons":["conf:1","win:tail","crowd:brand"],"computed_at":"2026-09-27T00:05:48Z"},"pay":null,"html_url":"https://alion.io/job/netapp-site-reliability-engineer-opensearch","json_url":"https://alion.io/job/netapp-site-reliability-engineer-opensearch.json","meta":{"generated_at":"2026-09-27T00:05:48Z","cache_seconds":300,"methodology":"https://alion.io/methodology","terms":"https://alion.io/terms","contact":"https://alion.io/contact","api":"https://alion.io/developers","usage":{"tier":"crawler","counted_by":"address","units_charged":1,"used_today":27,"day_limit":5000,"remaining_today":4973,"minute_limit":60,"resets_at":"2026-09-28T00:00:00Z"}}}