{"id":1573605,"url":"https://alion.io/job/cim-group-senior-data-scientist","title":"Senior Data Scientist","company":{"id":707107,"name":"CIM Group","domain":"cimgroup.com","url":"https://alion.io/company/cim-group","size_band":"501-1000","is_staffing_agency":false,"employer_type":"direct","is_intermediary":false,"listed_via":null,"ats_vendor":"Lever","truth_index":{"grade":"A","score":93,"open_postings":31,"ghost_share":0,"stale_share":0.097,"repost_share":0,"time_to_fill_p50_days":76,"computed_at":"2026-10-06T05:45:30Z"}},"role":"Data Science","role_family":"Data Science","seniority":"senior","employment_type":"full_time","work_mode":"on_site","remote_scope":null,"remote_scope_basis":null,"remote_working_hours":null,"hiring_geo_confidence":"structured","locations":["Los Angeles, United States"],"countries":["US"],"hiring_countries":[],"hiring_countries_total":0,"salary":null,"salary_estimate":{"min_usd":115000,"max_usd":224000,"period":"year","method":"role_seniority_country_remote_cell","sample_n":960},"experience_years_min":5,"visa_sponsorship":false,"relocation_package":false,"has_equity":false,"technologies":[{"name":"C#","optional":false},{"name":"C++","optional":false},{"name":"CI/CD","optional":false},{"name":"Docker","optional":false},{"name":"Java","optional":false},{"name":"Machine Learning","optional":false},{"name":"MATLAB","optional":false},{"name":"MLFlow","optional":false},{"name":"MS SQL","optional":false},{"name":"NumPy","optional":false},{"name":"Pandas","optional":false},{"name":"PostgreSQL","optional":false},{"name":"Python","optional":false},{"name":"PyTorch","optional":false},{"name":"PyTorch C++","optional":false},{"name":"R","optional":false},{"name":"Rest API","optional":false},{"name":"Scikit-learn","optional":false},{"name":"SQL","optional":false},{"name":"TensorFlow","optional":false},{"name":"TensorFlow C++","optional":false},{"name":"Amazon SageMaker","optional":true},{"name":"Azure","optional":true},{"name":"CatBoost","optional":true},{"name":"D3.js","optional":true},{"name":"Git","optional":true},{"name":"JavaScript","optional":true},{"name":"LightGBM","optional":true},{"name":"Matplotlib","optional":true},{"name":"NLP","optional":true},{"name":"Power BI","optional":true},{"name":"Reinforcement Learning","optional":true},{"name":"Seaborn","optional":true},{"name":"Tableau","optional":true},{"name":"Time Series Forecasting","optional":true},{"name":"Vertex AI","optional":true},{"name":"XGBoost","optional":true}],"status":"live","first_seen_at":"2026-09-30T23:57:23Z","employer_posted_date":"2026-09-30","last_verified_at":"2026-10-07T00:54:21Z","board_verified":true,"closed_at":null,"days_open":6,"trust":{"level":"ok","repost_count":null,"flags":[],"days_open":6},"description":"POSITION PURPOSE:\nThe Senior Data Scientist will be a pivotal member of our Enterprise Data Management team, playing a critical role in transforming raw, complex data into actionable intelligence that drives superior investment outcomes and operational efficiencies across the firm and its portfolio companies. The role is strategically positioned at the intersection of deep financial domain expertise and cutting-edge data science, requiring a proactive approach to problem-solving and the ability to translate complex technical findings into clear, strategic recommendations for investment teams and senior leadership. In this role, you will bridge the gap between advanced predictive modeling and production engineering, owning the full model lifecycle from feature engineering and offline experimentation to containerized deployment, CI/CD automation, and post-deployment monitoring. The work will directly contribute to, and be enabled by, the firm’s foundational data assets and policies.\nRESPONSIBILITIES:\nConduct in-depth analysis of large datasets to identify trends, patterns, and insights that can inform business strategy.\nTranslate complex findings into clear, actionable recommendations.\nEnd-to-end machine learning delivery, including architecting, training, evaluating, and deploying production-grade predictive models and statistical algorithms across vast, disparate datasets.\nDesign and implement sophisticated AI-powered algorithms and predictive models to continuously monitor, analyze, and optimize company performance across all three investment platforms.\nModel governance & monitoring; implement automated monitoring for production models to detect data drift, concept drift, feature skew, and latency degradation. Establish retraining triggers and model registry governance.\nOwn the deployment and operationalization of ML models using MLflow, containerization (Docker), and automated CI/CD pipelines. Transition models from notebook prototypes into scalable batch jobs or low-latency REST APIs.\nContinuously refine search parameters, algorithms, and recommendations based on historical deal patterns, dynamic market conditions, and emerging sectors.\nContribute to the design of scalable data pipelines and data lakes that power advanced analytics and AI applications.\nChampion best practices for data protection, including encryption and access controls, to safeguard sensitive information.\nCollaborate to define and implement robust data quality standards, ensuring all analytical models are built on a foundation of reliable data.\nOperate within our compliance framework to ensure ethical data handling, regulatory compliance, and consistency across the enterprise.\nWork with business stakeholders, data engineers, data stewards, and information architects to ensure data quality and accuracy of analytics and reporting.\nEDUCATION/EXPERIENCE REQUIREMENTS: (including certification, licenses, etc.)\nRequired:\nMaster’s or Bachelor’s degree in Data Science, Computer Science, Statistics, Mathematics, Quantitative Finance, or a closely related quantitative field.\n5+ years of dedicated experience as a Data Scientist.\nExpert-level Python and strong SQL skills, with experience using NumPy, Pandas, Scikit-learn, TensorFlow/PyTorch, and SQL databases such as PostgreSQL or T-SQL. (R, MATLAB, C++, Java, or C# are a plus.)\nDeep practical expertise with tree-based algorithms and gradient boosting (XGBoost, LightGBM, CatBoost) on structured/tabular business and financial data.\nStrong foundational statistics, including hypothesis testing, regression analysis, regularized models, classification, and time-series forecasting.\nStrong understanding of MLOps solutions for model deployment, management, and scaling (e.g., SageMaker, Vertex AI, Azure ML, MLflow, Docker).\nProven experience managing the full model lifecycle using MLflow, including experiment tracking, model registry, and artifact storage.\nWide range of ML algorithms, including k-NN, Naive Bayes, SVM, Decision Forests, Boosting Algorithms, Deep Learning, Time Series Forecasting, NLP, and Reinforcement Learning.\nAbility to work effectively in a dynamic, team-oriented environment.\nPreferred:\nRelevant professional certifications such as Chartered Financial Analyst (CFA) or Financial Risk Manager (FRM).\nStrong working knowledge of one of the following:Tableau, Power BI, D3.js, Matplotlib, Seaborn, or ggplot.\n\nFamiliarity with Git-based CI/CD workflows and automated testing for data science repositories.\nABOUT YOU:\nExcellent communication skills to articulate complex data concepts to non-technical stakeholders.\nAbility to build strong relationships across departments to ensure collaborative data initiatives.\nPragmatic problem solver; prioritize robust, explainable baseline models (like tuned gradient boosted trees) over unnecessary complexity, scaling up architecture only when the business problem demands it.\nStrong analytical and problem-solving skills to innovate and drive improvements in data processes are required.","description_format":"text","description_chars":5023,"description_truncated":false,"requirements":{"experience_years_min":5,"management_years_min":null,"team_size_min":null,"manages_managers":false,"education":{"level":"bachelor","optional":false},"security_clearance":false,"languages":[]},"benefits":[],"hiring_locations":[],"hiring_excludes":[],"relocation_offered":false,"industries":["Real Estate Investment","Property Development"],"lifecycle":[{"event":"open","at":"2026-10-01T08:58:45Z"}],"visa":[{"country":"US","licensed_sponsor":true,"evidence":"H-1B filings in 12 months: 6 · green card filings: 2","filings_12m":6,"filings_prev_12m":2,"green_card_filings_12m":2,"median_offered_wage_usd":190000,"route":null,"cap_exempt":false,"checked_at":"2026-10-03T21:08:04+00:00","sources":["US Department of Labor: LCA disclosure data (H-1B, H-1B1, E-3)","US Department of Labor: PERM disclosure data (green cards)"],"filings_for_role_12m":0}],"liveness":{"score":90,"band":"hot","label":"Hiring now","p_open":1,"p_active":0.903,"p_room":1,"age_days":5,"expected_fill_days":76,"reasons":["conf:2","velocity","win:early"],"computed_at":"2026-10-06T05:45:30Z"},"pay":null,"html_url":"https://alion.io/job/cim-group-senior-data-scientist","json_url":"https://alion.io/job/cim-group-senior-data-scientist.json","meta":{"generated_at":"2026-10-07T02:13:06Z","cache_seconds":300,"methodology":"https://alion.io/methodology","terms":"https://alion.io/terms","contact":"https://alion.io/contact","api":"https://alion.io/developers","about":"Alion is a live layer of people, companies and AI agents: who they are, whether they are real and active right now, what they do and how to work with them, readable by people and by agents and paid per call.","catalog":"https://alion.io/catalog.json","usage":{"tier":"crawler","counted_by":"address","units_charged":1,"used_today":3658,"day_limit":5000,"remaining_today":1342,"minute_limit":60,"resets_at":"2026-10-08T00:00:00Z"}},"offers":[{"id":"company.slices","title":"One company in depth, by slice","status":"live","price":{"credits":0.02,"usd":0.002,"plus_per_slice":{"credits":0.05,"usd":0.005}},"unit":"per company, plus each slice with data","note":"the employer in depth","call":{"mcp_tool":"get_company","arguments":{"id":707107},"rest":"https://alion.io/mcp/rest/get_company?id=707107"},"human":"https://alion.io/catalog?offer=company.slices&for=job%2Fcim-group-senior-data-scientist"},{"id":"market.stats","title":"A market slice: pay, demand and time to fill","status":"live","price":{"credits":1,"usd":0.1},"unit":"per slice","note":"pay, demand and time to fill for this role and place","call":{"mcp_tool":"market_stats"},"human":"https://alion.io/catalog?offer=market.stats&for=job%2Fcim-group-senior-data-scientist"},{"id":"job.search","title":"Open jobs by role, technology, place, pay and visa","status":"live","price":{"credits":0.02,"usd":0.002},"unit":"per posting in a list","note":"similar open postings","call":{"mcp_tool":"search_jobs"},"human":"https://alion.io/catalog?offer=job.search&for=job%2Fcim-group-senior-data-scientist"},{"id":"company.verify","title":"Is this company real and active right now","status":"pilot","price":null,"unit":"per company","request":{"url":"https://alion.io/catalog/request","method":"POST","body":"{\"offer\": \"company.verify\", \"for\": \"job/cim-group-senior-data-scientist\", \"note\": \"what you need it for\"}"},"human":"https://alion.io/catalog?offer=company.verify&for=job%2Fcim-group-senior-data-scientist"}]}