{"id":"d97110a2-c18b-4e1d-9261-79482f35eb45","entityType":"agent","slug":"crewai-geetanjalik01-ai-multi-agent-data-analysis-system","name":"AI-Multi-Agent-Data-Analysis-System","canonicalUrl":"https://www.xpersona.co/agent/crewai-geetanjalik01-ai-multi-agent-data-analysis-system","canonicalPath":"/agent/crewai-geetanjalik01-ai-multi-agent-data-analysis-system","generatedAt":"2026-10-09T22:49:59.210Z","source":"GITHUB_REPOS","claimStatus":"UNCLAIMED","verificationTier":"NONE","summary":{"evidence":{"source":"editorial-content","verified":true,"confidence":"high","updatedAt":"2026-10-09T15:57:33.992Z","emptyReason":null},"description":"AI-Multi-Agent-Data-Analysis-System An autonomous data engineering and predictive analytics pipeline via LangGraph and CrewAI, featuring structured data extraction, automated statistical cleaning, and machine learning inference. AI Multi-Agent Data Analysis System An AI-powered end-to-end data analysis platform that automates data cleaning, exploratory data analysis (EDA), visualization, machine learning, report generation, and semantic report search. Built using **Python**, **Streamlit**, **Scikit-learn**, **LangGraph**, **CrewAI**, **PostgreSQL**, and **ChromaDB**. --- Overview Analyzing datasets usually involves multiple manual steps such","descriptionLabel":"Technical summary","evidenceSummary":"Capability contract not published. No trust telemetry is available yet. Last updated 10/9/2026.","installCommand":null,"sourceUrl":"https://github.com/geetanjalik01/AI-Multi-Agent-Data-Analysis-System","homepage":null,"primaryLinks":[{"label":"View Source","url":"https://github.com/geetanjalik01/AI-Multi-Agent-Data-Analysis-System","kind":"source"}],"safetyScore":66,"overallRank":20,"popularityScore":0,"trustScore":null,"claimedByName":null,"isOwner":false,"seoDescription":"AI-Multi-Agent-Data-Analysis-System An autonomous data engineering and predictive analytics pipeline via LangGraph and CrewAI, featuring structured data extract"},"coverage":{"evidence":{"source":"public-profile","verified":false,"confidence":"medium","updatedAt":"2026-10-09T15:57:33.992Z","emptyReason":null},"protocols":[{"protocol":"OPENCLEW","label":"OpenClaw","status":"self-declared","notes":"Declared in the public agent profile."}],"capabilities":[{"label":"crewai","status":"self-declared"},{"label":"multi-agent","status":"self-declared"}],"verifiedCount":0,"selfDeclaredCount":3,"capabilityMatrix":{"rows":[{"key":"OPENCLEW","type":"protocol","support":"unknown","confidenceSource":"profile","notes":"Listed on profile"},{"key":"crewai","type":"capability","support":"supported","confidenceSource":"profile","notes":"Declared in agent profile metadata"},{"key":"multi-agent","type":"capability","support":"supported","confidenceSource":"profile","notes":"Declared in agent profile metadata"}],"flattenedTokens":"protocol:OPENCLEW|unknown|profile capability:crewai|supported|profile capability:multi-agent|supported|profile"}},"adoption":{"evidence":{"source":"no-adoption-signals","verified":false,"confidence":"low","updatedAt":"2026-10-09T15:57:33.992Z","emptyReason":"No source adoption metrics were available."},"stars":0,"forks":0,"downloads":null,"packageName":null,"latestVersion":null,"tractionLabel":null},"release":{"evidence":{"source":"agent-index","verified":false,"confidence":"medium","updatedAt":"2026-10-09T15:57:33.984Z","emptyReason":null},"lastUpdatedAt":"2026-10-09T15:57:33.992Z","lastCrawledAt":"2026-10-09T15:57:33.984Z","lastIndexedAt":null,"nextCrawlAt":"2026-10-16T15:57:33.984Z","lastVerifiedAt":null,"highlights":[]},"execution":{"evidence":{"source":"GITHUB REPOS","verified":false,"confidence":"low","updatedAt":null,"emptyReason":"No published capability contract is available yet."},"installCommand":null,"setupComplexity":"low","setupSteps":["Setup complexity is LOW. This package is likely designed for quick installation with minimal external side-effects.","Final validation: Expose the agent to a mock request payload inside a sandbox and trace the network egress before allowing access to real customer data."],"contract":{"contractStatus":"missing","authModes":[],"requires":[],"forbidden":[],"supportsMcp":false,"supportsA2a":false,"supportsStreaming":false,"inputSchemaRef":null,"outputSchemaRef":null,"dataRegion":null,"contractUpdatedAt":null,"sourceUpdatedAt":null,"freshnessSeconds":null},"invocationGuide":{"preferredApi":{"snapshotUrl":"https://www.xpersona.co/api/v1/agents/crewai-geetanjalik01-ai-multi-agent-data-analysis-system/snapshot","contractUrl":"https://www.xpersona.co/api/v1/agents/crewai-geetanjalik01-ai-multi-agent-data-analysis-system/contract","trustUrl":"https://www.xpersona.co/api/v1/agents/crewai-geetanjalik01-ai-multi-agent-data-analysis-system/trust"},"curlExamples":["curl -s \"https://www.xpersona.co/api/v1/agents/crewai-geetanjalik01-ai-multi-agent-data-analysis-system/snapshot\"","curl -s \"https://www.xpersona.co/api/v1/agents/crewai-geetanjalik01-ai-multi-agent-data-analysis-system/contract\"","curl -s \"https://www.xpersona.co/api/v1/agents/crewai-geetanjalik01-ai-multi-agent-data-analysis-system/trust\""],"jsonRequestTemplate":{"query":"summarize this repo","constraints":{"maxLatencyMs":2000,"protocolPreference":["OPENCLEW"]}},"jsonResponseTemplate":{"ok":true,"result":{"summary":"...","confidence":0.9},"meta":{"source":"GITHUB_REPOS","generatedAt":"2026-10-09T22:49:59.210Z"}},"retryPolicy":{"maxAttempts":3,"backoffMs":[500,1500,3500],"retryableConditions":["HTTP_429","HTTP_503","NETWORK_TIMEOUT"]}},"endpoints":{"dossierUrl":"https://www.xpersona.co/api/v1/agents/crewai-geetanjalik01-ai-multi-agent-data-analysis-system/dossier","snapshotUrl":"https://www.xpersona.co/api/v1/agents/crewai-geetanjalik01-ai-multi-agent-data-analysis-system/snapshot","contractUrl":"https://www.xpersona.co/api/v1/agents/crewai-geetanjalik01-ai-multi-agent-data-analysis-system/contract","trustUrl":"https://www.xpersona.co/api/v1/agents/crewai-geetanjalik01-ai-multi-agent-data-analysis-system/trust"}},"reliability":{"evidence":{"source":"runtime-metrics","verified":false,"confidence":"low","updatedAt":null,"emptyReason":"No trust, reliability, or runtime telemetry is available."},"trust":{"status":"unavailable","handshakeStatus":"UNKNOWN","verificationFreshnessHours":null,"reputationScore":null,"p95LatencyMs":null,"successRate30d":null,"fallbackRate":null,"attempts30d":null,"trustUpdatedAt":null,"trustConfidence":"unknown","sourceUpdatedAt":null,"freshnessSeconds":null},"decisionGuardrails":{"doNotUseIf":["Contract metadata is missing or unavailable for deterministic execution."],"safeUseWhen":[],"riskFlags":["missing_or_unavailable_contract","trust_data_unavailable","schema_references_missing"],"operationalConfidence":"low"},"executionMetrics":{"observedLatencyMsP50":null,"observedLatencyMsP95":null,"estimatedCostUsd":null,"uptime30d":null,"rateLimitRpm":null,"rateLimitBurst":null,"lastVerifiedAt":null,"verificationSource":null},"runtimeMetrics":{"successRate":null,"avgLatencyMs":null,"avgCostUsd":null,"hallucinationRate":null,"retryRate":null,"disputeRate":null,"p50Latency":null,"p95Latency":null,"lastUpdated":null}},"benchmarks":{"evidence":{"source":"no-benchmark-data","verified":false,"confidence":"low","updatedAt":null,"emptyReason":"No benchmark suites or observed failure patterns are available."},"suites":[],"failurePatterns":[]},"artifacts":{"evidence":{"source":"GITHUB REPOS","verified":false,"confidence":"high","updatedAt":"2026-10-09T15:57:33.992Z","emptyReason":null},"readme":"#  AI Multi-Agent Data Analysis System\n\nAn AI-powered end-to-end data analysis platform that automates data cleaning, exploratory data analysis (EDA), visualization, machine learning, report generation, and semantic report search.\n\nBuilt using **Python**, **Streamlit**, **Scikit-learn**, **LangGraph**, **CrewAI**, **PostgreSQL**, and **ChromaDB**.\n\n---\n\n## Overview\n\nAnalyzing datasets usually involves multiple manual steps such as cleaning data, performing exploratory analysis, generating visualizations, training machine learning models, and preparing reports.\n\nThis project automates the complete workflow. Users simply upload a CSV dataset, and the system generates meaningful insights, visualizations, machine learning predictions, and a business report automatically.\n\n---\n\n## Features\n\n-  Upload CSV datasets through a Streamlit interface\n-  Automatic data cleaning\n  - Remove duplicate records\n  - Handle missing values\n  - Standardize column names\n-  Exploratory Data Analysis (EDA)\n  - Dataset summary\n  - Descriptive statistics\n  - Correlation analysis\n-  Automatic visualizations\n  - Histograms\n  - Correlation Heatmaps\n-  Machine Learning\n  - Automatic Classification/Regression detection\n  - Decision Tree\n  - Random Forest\n  - Performance evaluation\n-  Automatic PDF report generation\n-  Store analysis history using PostgreSQL\n-  Semantic report search using ChromaDB + Sentence Transformers\n-  Modular AI Agent architecture using CrewAI\n-  Workflow organization using LangGraph\n\n---\n\n#  Workflow\n\n1. User uploads a CSV dataset.\n2. Dataset is cleaned automatically.\n3. Exploratory Data Analysis is performed.\n4. Visualizations are generated.\n5. ML module detects Classification or Regression automatically.\n6. Decision Tree and Random Forest models are trained.\n7. Best model performance is displayed.\n8. PDF business report is generated.\n9. Analysis history is stored in PostgreSQL.\n10. Report embeddings are stored in ChromaDB for semantic retrieval.\n\n---\n\n#  AI Components\n\n## CrewAI\n\nSpecialized AI agents were defined for:\n\n- Supervisor Agent\n- Cleaning Agent\n- EDA Agent\n- Visualization Agent\n- Machine Learning Agent\n- Report Generation Agent\n\nThese agents use **Llama 3.3 70B** through the **Groq API** for reasoning and modular workflow design.\n\n> **Note:** In the current implementation, the data processing pipeline is executed through Python modules, while CrewAI defines the modular multi-agent architecture.\n\n---\n\n## LangGraph\n\nLangGraph is used to organize the workflow by passing a shared workflow state between different modules.\n\n---\n\n## RAG Pipeline\n\nGenerated reports are:\n\n- Converted into embeddings using Sentence Transformers\n- Stored in ChromaDB\n- Retrieved using semantic similarity search\n\n---\n\n#  Tech Stack\n\n### Languages\n\n- Python\n\n### Data Analysis\n\n- Pandas\n- NumPy\n\n### Machine Learning\n\n- Scikit-learn\n- Decision Tree\n- Random Forest\n\n### Visualization\n\n- Matplotlib\n- Seaborn\n\n### Frontend\n\n- Streamlit\n\n### AI Frameworks\n\n- CrewAI\n- LangGraph\n\n### LLM\n\n- Llama 3.3 70B (Groq API)\n\n### Database\n\n- PostgreSQL\n- SQLAlchemy\n\n### Vector Database\n\n- ChromaDB\n\n### Embeddings\n\n- Sentence Transformers\n\n---\n\n#  Project Structure\n\n```\napp.py\nagents/\ntools/\ndatabase/\nrag/\nreports/\noutputs/\nuploads/\n```\n\n---\n\n# Future Improvements\n\n- Interactive Plotly dashboards\n- Hyperparameter tuning\n- Additional ML algorithms (XGBoost, LightGBM)\n- Multi-file support (Excel, JSON)\n- Cloud deployment (AWS/Azure)\n- Full CrewAI execution pipeline\n\n---\n\nB.Tech Information Technology | NIT Raipur\n\nInterested in Data Analytics, Machine Learning, AI, and Multi-Agent Systems.\n","readmeExcerpt":"AI Multi-Agent Data Analysis System An AI-powered end-to-end data analysis platform that automates data cleaning, exploratory data analysis (EDA), visualization, machine learning, report generation, and semantic report search. Built using **Python**, **Streamlit**, **Scikit-learn**, **LangGraph**, **CrewAI**, **PostgreSQL**, and **ChromaDB**. --- Overview Analyzing datasets usually involves multiple manual steps such","codeSnippets":[],"executableExamples":[{"language":"text","snippet":"app.py\nagents/\ntools/\ndatabase/\nrag/\nreports/\noutputs/\nuploads/"}],"parameters":null,"dependencies":[],"permissions":[],"extractedFiles":[],"languages":["python"],"docsSourceLabel":"GITHUB REPOS","editorialOverview":"AI-Multi-Agent-Data-Analysis-System An autonomous data engineering and predictive analytics pipeline via LangGraph and CrewAI, featuring structured data extraction, automated statistical cleaning, and machine learning inference. AI Multi-Agent Data Analysis System An AI-powered end-to-end data analysis platform that automates data cleaning, exploratory data analysis (EDA), visualization, machine learning, report generation, and semantic report search. Built using **Python**, **Streamlit**, **Scikit-learn**, **LangGraph**, **CrewAI**, **PostgreSQL**, and **ChromaDB**. --- Overview Analyzing datasets usually involves multiple manual steps such","editorialQuality":{"score":100,"threshold":65,"status":"ready","wordCount":428,"uniquenessScore":61,"reasons":[]}},"media":{"evidence":{"source":"no-media","verified":false,"confidence":"low","updatedAt":"2026-10-09T15:57:33.992Z","emptyReason":"No screenshots, media assets, or demo links are available."},"primaryImageUrl":null,"mediaAssetCount":0,"assets":[],"demoUrl":null},"ownerResources":{"evidence":{"source":"unclaimed","verified":false,"confidence":"low","updatedAt":"2026-10-09T15:57:33.992Z","emptyReason":"This page has not been claimed by the agent owner."},"hasCustomPage":false,"customPageUpdatedAt":null,"customLinks":[],"structuredLinks":{"docsUrl":null,"demoUrl":null,"supportUrl":null,"pricingUrl":null,"statusUrl":null},"customPage":null},"relatedAgents":{"evidence":{"source":"protocol-neighbors","verified":false,"confidence":"medium","updatedAt":"2026-10-09T22:49:59.210Z","emptyReason":null},"items":[{"id":"8ebccd8e-3863-4187-8355-c3f14e1f9edf","entityType":"agent","canonicalPath":"/agent/iofficeai-aionui","slug":"iofficeai-aionui","name":"AionUi","description":"Free, local, open-source 24/7 Cowork app and OpenClaw for Gemini CLI, Claude Code, Codex, OpenCode, Qwen Code, Goose CLI, Auggie, and more | 🌟 Star if you like it!","url":"https://github.com/iOfficeAI/AionUi","homepage":"https://www.aionui.com","source":"GITHUB_REPOS","protocols":["MCP","OPENCLAW"],"capabilities":[],"safetyScore":100,"overallRank":70,"updatedAt":"2026-10-09T19:11:12.944Z","createdAt":"2026-02-25T03:38:16.584Z","downloads":null},{"id":"b917f68a-ebff-438e-84f8-3f4b2494c0bc","entityType":"agent","canonicalPath":"/agent/activepieces-activepieces","slug":"activepieces-activepieces","name":"activepieces","description":"AI Agents & MCPs & AI Workflow Automation • (~400 MCP servers for AI agents) • AI Automation / AI Agent with MCPs • AI Workflows & AI Agents • MCPs for AI Agents","url":"https://github.com/activepieces/activepieces","homepage":"https://www.activepieces.com","source":"GITHUB_REPOS","protocols":["OPENCLAW"],"capabilities":[],"safetyScore":100,"overallRank":70,"updatedAt":"2026-04-15T02:22:12.426Z","createdAt":"2026-02-25T03:38:12.412Z","downloads":null},{"id":"5cb26759-3a39-483f-94cf-276a98c13bb8","entityType":"agent","canonicalPath":"/agent/cherryhq-cherry-studio","slug":"cherryhq-cherry-studio","name":"cherry-studio","description":"AI productivity studio with smart chat, autonomous agents, and 300+ assistants. Unified access to frontier LLMs","url":"https://github.com/CherryHQ/cherry-studio","homepage":"https://cherry-ai.com","source":"GITHUB_REPOS","protocols":["MCP","OPENCLAW"],"capabilities":[],"safetyScore":100,"overallRank":70,"updatedAt":"2026-04-11T14:38:40.986Z","createdAt":"2026-02-25T03:38:19.379Z","downloads":null},{"id":"6f6582d0-5d76-4f0f-b81d-86520247950b","entityType":"agent","canonicalPath":"/agent/copilotkit-copilotkit","slug":"copilotkit-copilotkit","name":"CopilotKit","description":"The Frontend for Agents & Generative UI. React + Angular","url":"https://github.com/CopilotKit/CopilotKit","homepage":"https://docs.copilotkit.ai","source":"GITHUB_REPOS","protocols":["OPENCLAW"],"capabilities":[],"safetyScore":100,"overallRank":70,"updatedAt":"2026-03-25T09:50:57.846Z","createdAt":"2026-02-25T03:39:14.617Z","downloads":null}],"links":{"hub":"/agent","source":"/agent/source/github_repos","protocols":[{"label":"OpenClaw","href":"/agent/protocol/openclew"}]}}}