{"adoption": {"forks": 65, "observed_at": "2026-08-28T04:04:08.818012+00:00", "stars": 1254}, "canonical_url": "https://ross.abutalabs.com/products/llamagym", "card": {"archived": false, "artifact_type": "library", "description": "Fine-tune LLM agents with online reinforcement learning", "domain": ["reinforcement-learning", "large-language-models", "machine-learning"], "enriched": true, "function": ["reinforcement-learning", "llm-training", "agent-framework", "machine-learning"], "health_score": 20, "homepage": null, "language": "Python", "license": "MIT", "license_family": "permissive", "maturity": "experimental", "member_repos": ["KhoomeiK/LlamaGym"], "name": "KhoomeiK/LlamaGym", "platform": ["python"], "pushed_at": "2024-03-19T17:34:28+00:00", "repo": "KhoomeiK/LlamaGym", "stars": 1254, "tags": ["gym-environments", "fine-tuning", "ppo", "abstract-class", "openai-gym", "ai-agents"], "topics": [], "urls": [], "use_cases": ["fine-tune an LLM agent with reinforcement learning in a Gym environment", "train an LLM to play blackjack via RL", "experiment with agent prompts and hyperparameters across RL environments", "simplify PPO setup for LLM agents", "build agents that learn online from reward signals"], "what_it_is": "LlamaGym is a Python library that simplifies fine-tuning LLM-based agents with online reinforcement learning in Gym-style environments. It provides a single Agent abstract class that handles conversation context, episode batching, reward assignment, and PPO setup.", "when_to_avoid": ["you need a production-ready, actively maintained RL training framework", "you want offline fine-tuning without reinforcement learning", "you need multi-agent or non-Gym training setups"], "when_to_choose": ["you want to fine-tune an LLM agent with RL in a Gymnasium environment without writing boilerplate", "you want a minimal abstract class to iterate on agent prompting and hyperparameters", "you already have a Gym environment and an LLM you want to train"]}, "data_as_of": "2026-08-30T08:39:29.467469+00:00", "members": [{"path": "/products/llamagym", "repo": "KhoomeiK/LlamaGym", "role": "main", "score": 25}], "provenance": {"archived": {"kind": "observed", "observed_at": "2026-08-28T04:04:08.818012+00:00", "source": "github"}, "artifact_type": {"confidence": null, "enriched_at": "2026-08-30T05:07:23.092990+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "36ebe41314d645afcc98e70cbe5cad770e2a988f4d83fabdbbeb76a6addae9b1", "fetched_at": "2026-08-28T04:04:08.818012+00:00", "kind": "readme", "missing": false, "url": "https://github.com/KhoomeiK/LlamaGym"}], "taxonomy_version": 1}, "description": {"kind": "observed", "observed_at": "2026-08-28T04:04:08.818012+00:00", "source": "github"}, "domain": {"confidence": null, "enriched_at": "2026-08-30T05:07:23.092990+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "36ebe41314d645afcc98e70cbe5cad770e2a988f4d83fabdbbeb76a6addae9b1", "fetched_at": "2026-08-28T04:04:08.818012+00:00", "kind": "readme", "missing": false, "url": "https://github.com/KhoomeiK/LlamaGym"}], "taxonomy_version": 1}, "enriched": {"inputs": [], "kind": "computed", "method": "enrichment_status"}, "function": {"confidence": null, "enriched_at": "2026-08-30T05:07:23.092990+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "36ebe41314d645afcc98e70cbe5cad770e2a988f4d83fabdbbeb76a6addae9b1", "fetched_at": "2026-08-28T04:04:08.818012+00:00", "kind": "readme", "missing": false, "url": "https://github.com/KhoomeiK/LlamaGym"}], "taxonomy_version": 1}, "health_score": {"inputs": ["days_since_push", "days_since_release", "archived"], "kind": "computed", "method": "health_v1"}, "homepage": {"kind": "observed", "observed_at": "2026-08-28T04:04:08.818012+00:00", "source": "github"}, "language": {"kind": "observed", "observed_at": "2026-08-28T04:04:08.818012+00:00", "source": "github"}, "license": {"kind": "observed", "observed_at": "2026-08-28T04:04:08.818012+00:00", "source": "github"}, "license_family": {"inputs": ["license"], "kind": "computed", "method": "license_family"}, "maturity": {"confidence": null, "enriched_at": "2026-08-30T05:07:23.092990+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "36ebe41314d645afcc98e70cbe5cad770e2a988f4d83fabdbbeb76a6addae9b1", "fetched_at": "2026-08-28T04:04:08.818012+00:00", "kind": "readme", "missing": false, "url": "https://github.com/KhoomeiK/LlamaGym"}], "taxonomy_version": 1}, "member_repos": {"kind": "observed", "observed_at": "2026-08-28T04:04:08.818012+00:00", "source": "github"}, "name": {"kind": "observed", "observed_at": "2026-08-28T04:04:08.818012+00:00", "source": "github"}, "platform": {"confidence": null, "enriched_at": "2026-08-30T05:07:23.092990+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "36ebe41314d645afcc98e70cbe5cad770e2a988f4d83fabdbbeb76a6addae9b1", "fetched_at": "2026-08-28T04:04:08.818012+00:00", "kind": "readme", "missing": false, "url": "https://github.com/KhoomeiK/LlamaGym"}], "taxonomy_version": 1}, "pushed_at": {"kind": "observed", "observed_at": "2026-08-28T04:04:08.818012+00:00", "source": "github"}, "repo": {"kind": "observed", "observed_at": "2026-08-28T04:04:08.818012+00:00", "source": "github"}, "stars": {"kind": "observed", "observed_at": "2026-08-28T04:04:08.818012+00:00", "source": "github"}, "tags": {"confidence": null, "enriched_at": "2026-08-30T05:07:23.092990+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "36ebe41314d645afcc98e70cbe5cad770e2a988f4d83fabdbbeb76a6addae9b1", "fetched_at": "2026-08-28T04:04:08.818012+00:00", "kind": "readme", "missing": false, "url": "https://github.com/KhoomeiK/LlamaGym"}], "taxonomy_version": 1}, "topics": {"kind": "observed", "observed_at": "2026-08-28T04:04:08.818012+00:00", "source": "github"}, "urls": {"kind": "observed", "observed_at": "2026-08-28T04:04:08.818012+00:00", "source": "github"}, "use_cases": {"confidence": null, "enriched_at": "2026-08-30T05:07:23.092990+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "36ebe41314d645afcc98e70cbe5cad770e2a988f4d83fabdbbeb76a6addae9b1", "fetched_at": "2026-08-28T04:04:08.818012+00:00", "kind": "readme", "missing": false, "url": "https://github.com/KhoomeiK/LlamaGym"}], "taxonomy_version": 1}, "what_it_is": {"confidence": null, "enriched_at": "2026-08-30T05:07:23.092990+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "36ebe41314d645afcc98e70cbe5cad770e2a988f4d83fabdbbeb76a6addae9b1", "fetched_at": "2026-08-28T04:04:08.818012+00:00", "kind": "readme", "missing": false, "url": "https://github.com/KhoomeiK/LlamaGym"}], "taxonomy_version": 1}, "when_to_avoid": {"confidence": null, "enriched_at": "2026-08-30T05:07:23.092990+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "36ebe41314d645afcc98e70cbe5cad770e2a988f4d83fabdbbeb76a6addae9b1", "fetched_at": "2026-08-28T04:04:08.818012+00:00", "kind": "readme", "missing": false, "url": "https://github.com/KhoomeiK/LlamaGym"}], "taxonomy_version": 1}, "when_to_choose": {"confidence": null, "enriched_at": "2026-08-30T05:07:23.092990+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "36ebe41314d645afcc98e70cbe5cad770e2a988f4d83fabdbbeb76a6addae9b1", "fetched_at": "2026-08-28T04:04:08.818012+00:00", "kind": "readme", "missing": false, "url": "https://github.com/KhoomeiK/LlamaGym"}], "taxonomy_version": 1}}, "score": {"components": {"activity": 0, "longevity": 65, "rhythm": 35}, "computed_at": "2026-09-02T17:46:02.011165+00:00", "flags": ["no_releases"], "formula": "round(0.45*activity + 0.35*rhythm + 0.20*longevity); archived -> min(score, 10)", "inputs": {"age_days": 915, "days_push": 897, "days_rel": null, "gap_med": null, "n_releases_24m": 0}, "score": 25, "version": 2}, "staleness": {"enrichment_outdated": false, "low_confidence": false, "scrape_days": 9, "stale_scrape": false}}