{"adoption": {"forks": 246, "observed_at": "2026-08-28T04:05:37.246377+00:00", "stars": 1793}, "canonical_url": "https://ross.abutalabs.com/products/datagen", "card": {"archived": false, "artifact_type": "application", "description": "DATAGEN: AI-driven multi-agent research assistant automating hypothesis generation, data analysis, and report writing. ", "domain": ["artificial-intelligence", "data-science", "large-language-models", "data-visualization", "analytics"], "enriched": true, "function": ["agent-framework", "data-science", "data-visualization", "llm-inference", "chatbot"], "health_score": 79, "homepage": null, "language": "Python", "license": "MIT", "license_family": "permissive", "maturity": "active", "member_repos": ["zi-yue-1129/DATAGEN"], "name": "zi-yue-1129/DATAGEN", "platform": ["python", "cross-platform"], "pushed_at": "2026-08-16T08:38:07+00:00", "repo": "zi-yue-1129/DATAGEN", "stars": 1793, "tags": ["multi-agent", "langchain", "langgraph", "hypothesis-generation", "automated-reporting", "research-assistant", "openai", "ai-agents"], "topics": ["artificial-intelligence", "data-analysis", "data-analytics", "data-science", "langchain", "langgraph", "large-language-model", "large-language-models", "multiagent-systems", "python", "ai-data-analysis", "ai", "code-generation", "agent", "llm"], "urls": [], "use_cases": ["automate exploratory data analysis with AI agents", "generate research hypotheses from a dataset", "produce analysis reports and charts automatically", "run a multi-agent data science pipeline", "analyze csv data and get insights without coding", "build an AI research assistant workflow"], "what_it_is": "DATAGEN is an AI-powered multi-agent research and data analysis platform built with LangChain, LangGraph, and OpenAI GPT models. It automates hypothesis generation, data cleaning and analysis, visualization, and report writing through coordinated specialized agents.", "when_to_avoid": ["you need fully deterministic, auditable analysis without LLM involvement", "you cannot share your data with external LLM APIs", "you need a lightweight single-script analysis tool", "you require strict enterprise compliance guarantees out of the box"], "when_to_choose": ["you want end-to-end automated data analysis and report generation", "you need hypothesis generation and validation driven by LLMs", "you want a Python-based multi-agent system built on LangGraph", "you prefer a self-hosted research assistant you can customize"]}, "data_as_of": "2026-08-30T08:39:29.467469+00:00", "members": [{"path": "/products/datagen", "repo": "zi-yue-1129/DATAGEN", "role": "main", "score": 67}], "provenance": {"archived": {"kind": "observed", "observed_at": "2026-08-28T04:05:37.246377+00:00", "source": "github"}, "artifact_type": {"confidence": null, "enriched_at": "2026-08-30T03:22:57.067580+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "4e98b39bae4211565e911c09644f7c8c70c488c3e7dc51c42affb0ff8b747a08", "fetched_at": "2026-08-28T04:05:37.246377+00:00", "kind": "readme", "missing": false, "url": "https://github.com/zi-yue-1129/DATAGEN"}], "taxonomy_version": 1}, "description": {"kind": "observed", "observed_at": "2026-08-28T04:05:37.246377+00:00", "source": "github"}, "domain": {"confidence": null, "enriched_at": "2026-08-30T03:22:57.067580+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "4e98b39bae4211565e911c09644f7c8c70c488c3e7dc51c42affb0ff8b747a08", "fetched_at": "2026-08-28T04:05:37.246377+00:00", "kind": "readme", "missing": false, "url": "https://github.com/zi-yue-1129/DATAGEN"}], "taxonomy_version": 1}, "enriched": {"inputs": [], "kind": "computed", "method": "enrichment_status"}, "function": {"confidence": null, "enriched_at": "2026-08-30T03:22:57.067580+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "4e98b39bae4211565e911c09644f7c8c70c488c3e7dc51c42affb0ff8b747a08", "fetched_at": "2026-08-28T04:05:37.246377+00:00", "kind": "readme", "missing": false, "url": "https://github.com/zi-yue-1129/DATAGEN"}], "taxonomy_version": 1}, "health_score": {"inputs": ["days_since_push", "days_since_release", "archived"], "kind": "computed", "method": "health_v1"}, "homepage": {"kind": "observed", "observed_at": "2026-08-28T04:05:37.246377+00:00", "source": "github"}, "language": {"kind": "observed", "observed_at": "2026-08-28T04:05:37.246377+00:00", "source": "github"}, "license": {"kind": "observed", "observed_at": "2026-08-28T04:05:37.246377+00:00", "source": "github"}, "license_family": {"inputs": ["license"], "kind": "computed", "method": "license_family"}, "maturity": {"confidence": null, "enriched_at": "2026-08-30T03:22:57.067580+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "4e98b39bae4211565e911c09644f7c8c70c488c3e7dc51c42affb0ff8b747a08", "fetched_at": "2026-08-28T04:05:37.246377+00:00", "kind": "readme", "missing": false, "url": "https://github.com/zi-yue-1129/DATAGEN"}], "taxonomy_version": 1}, "member_repos": {"kind": "observed", "observed_at": "2026-08-28T04:05:37.246377+00:00", "source": "github"}, "name": {"kind": "observed", "observed_at": "2026-08-28T04:05:37.246377+00:00", "source": "github"}, "platform": {"confidence": null, "enriched_at": "2026-08-30T03:22:57.067580+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "4e98b39bae4211565e911c09644f7c8c70c488c3e7dc51c42affb0ff8b747a08", "fetched_at": "2026-08-28T04:05:37.246377+00:00", "kind": "readme", "missing": false, "url": "https://github.com/zi-yue-1129/DATAGEN"}], "taxonomy_version": 1}, "pushed_at": {"kind": "observed", "observed_at": "2026-08-28T04:05:37.246377+00:00", "source": "github"}, "repo": {"kind": "observed", "observed_at": "2026-08-28T04:05:37.246377+00:00", "source": "github"}, "stars": {"kind": "observed", "observed_at": "2026-08-28T04:05:37.246377+00:00", "source": "github"}, "tags": {"confidence": null, "enriched_at": "2026-08-30T03:22:57.067580+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "4e98b39bae4211565e911c09644f7c8c70c488c3e7dc51c42affb0ff8b747a08", "fetched_at": "2026-08-28T04:05:37.246377+00:00", "kind": "readme", "missing": false, "url": "https://github.com/zi-yue-1129/DATAGEN"}], "taxonomy_version": 1}, "topics": {"kind": "observed", "observed_at": "2026-08-28T04:05:37.246377+00:00", "source": "github"}, "urls": {"kind": "observed", "observed_at": "2026-08-28T04:05:37.246377+00:00", "source": "github"}, "use_cases": {"confidence": null, "enriched_at": "2026-08-30T03:22:57.067580+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "4e98b39bae4211565e911c09644f7c8c70c488c3e7dc51c42affb0ff8b747a08", "fetched_at": "2026-08-28T04:05:37.246377+00:00", "kind": "readme", "missing": false, "url": "https://github.com/zi-yue-1129/DATAGEN"}], "taxonomy_version": 1}, "what_it_is": {"confidence": null, "enriched_at": "2026-08-30T03:22:57.067580+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "4e98b39bae4211565e911c09644f7c8c70c488c3e7dc51c42affb0ff8b747a08", "fetched_at": "2026-08-28T04:05:37.246377+00:00", "kind": "readme", "missing": false, "url": "https://github.com/zi-yue-1129/DATAGEN"}], "taxonomy_version": 1}, "when_to_avoid": {"confidence": null, "enriched_at": "2026-08-30T03:22:57.067580+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "4e98b39bae4211565e911c09644f7c8c70c488c3e7dc51c42affb0ff8b747a08", "fetched_at": "2026-08-28T04:05:37.246377+00:00", "kind": "readme", "missing": false, "url": "https://github.com/zi-yue-1129/DATAGEN"}], "taxonomy_version": 1}, "when_to_choose": {"confidence": null, "enriched_at": "2026-08-30T03:22:57.067580+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "4e98b39bae4211565e911c09644f7c8c70c488c3e7dc51c42affb0ff8b747a08", "fetched_at": "2026-08-28T04:05:37.246377+00:00", "kind": "readme", "missing": false, "url": "https://github.com/zi-yue-1129/DATAGEN"}], "taxonomy_version": 1}}, "score": {"components": {"activity": 98, "longevity": 55, "rhythm": 35}, "computed_at": "2026-09-02T17:46:02.011165+00:00", "flags": ["no_releases"], "formula": "round(0.45*activity + 0.35*rhythm + 0.20*longevity); archived -> min(score, 10)", "inputs": {"age_days": 771, "days_push": 17, "days_rel": null, "gap_med": null, "n_releases_24m": 0}, "score": 67, "version": 2}, "staleness": {"enrichment_outdated": false, "low_confidence": false, "scrape_days": 9, "stale_scrape": false}}