{"id":1247949,"url":"https://alion.io/job/beam-applied-ai-research-engineer","title":"Applied AI Research Engineer","company":{"id":1035,"name":"Beam","domain":"beam.cloud","url":"https://alion.io/company/beam","size_band":"51-200","is_staffing_agency":false,"employer_type":"direct","is_intermediary":false,"listed_via":null,"ats_vendor":"Work at a Startup","truth_index":null},"role":"AI/ML","role_family":"AI/ML","seniority":"middle","employment_type":"full_time","work_mode":"remote","remote_scope":"stated_countries","remote_scope_basis":"posting_text","remote_working_hours":null,"hiring_geo_confidence":"structured","locations":["New York, United States"],"countries":["US"],"hiring_countries":["US"],"hiring_countries_total":1,"salary":{"min":140000,"max":200000,"currency":"USD","period":"year","gross":null,"usd_annual":200000},"salary_estimate":null,"experience_years_min":3,"visa_sponsorship":false,"relocation_package":false,"has_equity":true,"technologies":[{"name":"GitHub","optional":false},{"name":"KV Cache","optional":false},{"name":"LLM","optional":false},{"name":"PyTorch","optional":false},{"name":"Quantization","optional":false},{"name":"Reinforcement Learning","optional":false},{"name":"Snyk","optional":false},{"name":"Speculative Decoding","optional":false}],"status":"live","first_seen_at":"2026-09-25T18:06:24Z","employer_posted_date":"2026-09-25","last_verified_at":"2026-09-27T00:25:21Z","board_verified":true,"closed_at":null,"days_open":1,"trust":{"level":"ok","repost_count":null,"flags":[],"days_open":1},"description":"Beam is an ultrafast AI inference platform. We built a serverless runtime that launches GPU-backed containers in less than 1 second and quickly scales out to thousands of GPUs. Developers use our platform to serve apps to millions of users around the globe. We're backed by Y Combinator, Tiger Global, and prominent developer-tool founders, including the founder of Snyk and former CTO of GitHub.\nAbout the Role\nWe’re looking to hire someone to own inference research hands-on and find ways to lower cost per token and latency on our customer workloads.\nLow-level inference optimization, from speculative decoding, quantization, KV-cache and memory management\nWork directly with customers to optimize their production workloads, and apply your learnings to our platform as product improvements\nHigh-level of autonomy to find the highest upside bets and guide the future of our inference platform based on your work\nSkills & Experience\nSystems or research background in LLM inference\nDeep understanding of LLM serving, from the kernel to the scheduler\nHistory of shipping products or research that people use in production-like scenarios, whether academic or industry\nExcited to collaborate closely with customers\nEnthusiasm for developer tools, cloud native technologies, and open source software\nBenefits\nCompetitive salary and meaningful equity\nJoin a fast-growing pre-series A company at the ground floor\nHealth, dental, and vision benefits with 90% coverage for you and 50% for dependents\nOpportunities to participate in events across the cloud native community\nFitness stipend, learning budget, and much, much more","description_format":"text","description_chars":1619,"description_truncated":false,"requirements":{"experience_years_min":3,"management_years_min":null,"team_size_min":null,"manages_managers":false,"education":null,"security_clearance":false,"languages":[]},"benefits":["Equity"],"hiring_locations":[{"name":"United States","iso":"US","kind":"country"}],"hiring_excludes":[],"relocation_offered":false,"industries":["Cloud Platforms (IaaS & PaaS)","AI Compute & Inference"],"lifecycle":[{"event":"open","at":"2026-09-25T18:06:24Z"}],"liveness":{"score":86,"band":"hot","label":"Hiring now","p_open":1,"p_active":0.86,"p_room":1,"age_days":0,"expected_fill_days":64,"reasons":["conf:4","win:early"],"computed_at":"2026-09-26T05:45:00Z"},"pay":{"stated_usd_annual":200000,"is_top_pay":true},"html_url":"https://alion.io/job/beam-applied-ai-research-engineer","json_url":"https://alion.io/job/beam-applied-ai-research-engineer.json","meta":{"generated_at":"2026-09-27T00:38:00Z","cache_seconds":300,"methodology":"https://alion.io/methodology","terms":"https://alion.io/terms","contact":"https://alion.io/contact","api":"https://alion.io/developers","usage":{"tier":"crawler","counted_by":"address","units_charged":1,"used_today":542,"day_limit":5000,"remaining_today":4458,"minute_limit":60,"resets_at":"2026-09-28T00:00:00Z"}}}