{"id":1273196,"url":"https://alion.io/job/slo-technologies-senior-data-analyst","title":"Senior Data Analyst","company":{"id":3801243,"name":"SLO Technologies","domain":"advarisk.com","url":"https://alion.io/company/slo-technologies","size_band":null,"is_staffing_agency":false,"employer_type":"direct","is_intermediary":false,"listed_via":null,"ats_vendor":null,"truth_index":null},"role":"Analytics","role_family":"Analytics","seniority":"senior","employment_type":null,"work_mode":"on_site","remote_scope":null,"remote_scope_basis":null,"remote_working_hours":null,"hiring_geo_confidence":"structured","locations":["Pune, India"],"countries":["IN"],"hiring_countries":[],"hiring_countries_total":0,"salary":null,"salary_estimate":{"min_usd":23000,"max_usd":51000,"period":"year","method":"global_role_cell_scaled_by_country","sample_n":1330},"experience_years_min":3,"visa_sponsorship":false,"relocation_package":false,"has_equity":false,"technologies":[{"name":"AI Agents","optional":false},{"name":"Hugging Face","optional":false},{"name":"Machine Learning","optional":false},{"name":"NER","optional":false},{"name":"NLP","optional":false},{"name":"NLTK","optional":false},{"name":"NumPy","optional":false},{"name":"OpenAI","optional":false},{"name":"Pandas","optional":false},{"name":"Prompt Engineering","optional":false},{"name":"Python","optional":false},{"name":"Rest API","optional":false},{"name":"Sentence-Transformers","optional":false},{"name":"spaCy","optional":false},{"name":"SQL","optional":false},{"name":"Transformers","optional":false},{"name":"Vertex AI","optional":false},{"name":"AWS","optional":true},{"name":"Azure","optional":true},{"name":"Embeddings","optional":true},{"name":"GCP","optional":true},{"name":"LangChain","optional":true},{"name":"LlamaIndex","optional":true},{"name":"LLM","optional":true},{"name":"RAG","optional":true}],"status":"live","first_seen_at":"2026-08-27T10:23:04Z","employer_posted_date":null,"last_verified_at":"2026-08-27T10:23:04Z","board_verified":false,"closed_at":null,"days_open":31,"trust":{"level":"not_scored","repost_count":null,"flags":[],"days_open":31},"description":"Key Responsibilities :\n\n- Analyze large volumes of structured and unstructured data to derive actionable insights.\n\n- Design, develop, and optimize Python scripts for data extraction, cleaning, transformation, and analysis of regional data.\n\n- Prepare, preprocess, and annotate datasets for NLP and machine learning applications.\n\nTrain, fine-tune, and evaluate NLP models, including :\n\n- Named Entity Recognition (NER)\n\n- Transliteration models\n\n- Text classification\n\n- Name Matching\n\n- Perform pattern identification and data mining to uncover trends, anomalies, and business insights.\n\n- Develop data quality frameworks and implement validation processes to improve model performance.\n\nCalculate and monitor model evaluation metrics, including :\n\n- Accuracy, Precision, Recall, F1-Score\n\n- Collaborate with AI/ML teams to integrate AI agents and automate business workflows.\n\n- Support prompt engineering, testing, and validation of AI-powered applications.\n\n- Work with cross-functional teams to understand business requirements and translate them into analytical solutions.\n\n- Document methodologies, workflows, datasets, and model performance reports.\n\n- Continuously improve data processing pipelines and contribute to AI solution enhancements.\n\nRequired Skills & Qualifications :\n\n- Bachelor's or Master's degree in Computer Science, Data Science, Artificial Intelligence, Statistics, or a related field.\n\n- Strong proficiency in Python for data analysis and automation.\n\nHands-on experience with NLP libraries and frameworks such as :\n\n- spaCy, NLTK, Hugging Face Transformers, Sentence Transformers\n\n- Experience training and evaluating NLP models.\n\n- Strong understanding of unstructured data processing techniques.\n\n- Experience in feature engineering, data preprocessing, and text normalization.\n\n- Good understanding of machine learning model evaluation metrics, including Accuracy, Precision, Recall, F1-Score, and Confusion Matrix.\n\n- Strong Experience with SQL and data manipulation libraries i.e. Pandas and NumPy.\n\n- Familiarity with REST APIs and AI agent integration frameworks ie. VertexAI / OpenAI\n\n- Strong analytical, problem-solving, and debugging skills.\n\n- Excellent communication and documentation abilities.\n\n- Mentoring to junior team members.\n\nPreferred Skills :\n\n- Experience with Large Language Models (LLMs) and Generative AI.\n\n- Knowledge of Retrieval-Augmented Generation (RAG) architectures.\n\n- Experience with vector databases and embeddings.\n\n- Familiarity with LangChain, LlamaIndex, or similar AI orchestration frameworks.\n\n- Experience with cloud platforms such as AWS, Azure, or Google Cloud.\n\n- Knowledge of MLOps concepts and model deployment pipelines.\n\n- Experience working with multilingual datasets.\n\n- Problem-solving mindset, Data-driven decision making, Team collaboration, Ownership\n\nand accountability\n\nExperience\n\n- 3-6 years of relevant experience in Data Analytics, NLP, AI, or Machine Learning.\nSkills\nPython, SQL, Pandas, Numpy, REST API, Data Analytics, Data Analyst, LLM, Machine Learning","description_format":"text","description_chars":3051,"description_truncated":false,"requirements":{"experience_years_min":3,"management_years_min":null,"team_size_min":null,"manages_managers":false,"education":null,"security_clearance":false,"languages":[]},"benefits":[],"hiring_locations":[],"hiring_excludes":[],"relocation_offered":false,"industries":["Financial Software & Embedded Finance"],"lifecycle":[{"event":"open","at":"2026-09-26T00:07:01Z"}],"liveness":{"score":21,"band":"cold","label":"Long shot","p_open":0.6,"p_active":0.625,"p_room":0.55,"age_days":30,"expected_fill_days":24,"reasons":["seen:30","win:tail"],"computed_at":"2026-09-27T05:45:00Z"},"pay":null,"html_url":"https://alion.io/job/slo-technologies-senior-data-analyst","json_url":"https://alion.io/job/slo-technologies-senior-data-analyst.json","meta":{"generated_at":"2026-09-28T03:41:14Z","cache_seconds":300,"methodology":"https://alion.io/methodology","terms":"https://alion.io/terms","contact":"https://alion.io/contact","api":"https://alion.io/developers","usage":{"tier":"crawler","counted_by":"address","units_charged":1,"used_today":2292,"day_limit":5000,"remaining_today":2708,"minute_limit":60,"resets_at":"2026-09-29T00:00:00Z"}}}