{"adoption": {"forks": 1005, "observed_at": "2026-08-28T04:10:38.955510+00:00", "stars": 9956}, "canonical_url": "https://ross.abutalabs.com/products/openrlhf", "card": {"archived": false, "artifact_type": "framework", "description": "An Easy-to-use, Scalable and High-performance Agentic RL Framework based on Ray (PPO & DAPO & REINFORCE++ &  VLM & TIS & vLLM & Ray & Async  RL)", "domain": ["reinforcement-learning", "large-language-models", "machine-learning", "gpu-computing"], "enriched": true, "function": ["llm-training", "reinforcement-learning", "machine-learning", "gpu-computing", "cli"], "health_score": 98, "homepage": "https://openrlhf.readthedocs.io/", "language": "Python", "license": "Apache-2.0", "license_family": "permissive", "maturity": "active", "member_repos": ["OpenRLHF/OpenRLHF"], "name": "OpenRLHF/OpenRLHF", "platform": ["python"], "pushed_at": "2026-08-13T11:26:24+00:00", "repo": "OpenRLHF/OpenRLHF", "stars": 9956, "tags": ["rlhf", "ppo", "grpo", "vllm", "ray", "deepspeed", "dpo", "vlm-training", "distributed-training", "agent-rl", "linux", "gpu", "docker"], "topics": ["transformers", "vllm", "large-language-models", "raylib", "reinforcement-learning-from-human-feedback", "reinforcement-learning", "proximal-policy-optimization", "visual-language-models"], "urls": [], "use_cases": ["train llm with rlhf ppo", "run grpo reinforcement learning fine-tuning", "fine-tune a 70b model with reinforcement learning from human feedback", "train vision-language model with rl", "sft and reward model training pipeline", "async rl training with vllm rollout", "multi-turn agent reinforcement learning training"], "what_it_is": "OpenRLHF is a high-performance, production-ready open-source RLHF framework built on Ray + vLLM + DeepSpeed for scalable reinforcement learning from human feedback. It supports state-of-the-art RL algorithms (PPO, GRPO, REINFORCE++, RLOO), SFT, reward modeling, DPO, and vision-language model training through a unified agent-based pipeline.", "when_to_avoid": ["you only need simple supervised fine-tuning without RL", "you lack multi-GPU infrastructure or distributed training experience", "you need a lightweight single-GPU hobbyist trainer", "you prefer a non-Python training stack"], "when_to_choose": ["you need scalable RLHF/RL post-training for large language models on multi-GPU clusters", "you want to switch between PPO, GRPO, REINFORCE++, and RLOO with a single flag", "you need vLLM-accelerated generation and DeepSpeed ZeRO-3 training from HuggingFace checkpoints", "you want to train VLMs or multi-turn agents end-to-end with RL"]}, "data_as_of": "2026-08-30T08:39:29.467469+00:00", "members": [{"path": "/products/openrlhf", "repo": "OpenRLHF/OpenRLHF", "role": "main", "score": 89}], "provenance": {"archived": {"kind": "observed", "observed_at": "2026-08-28T04:10:38.955510+00:00", "source": "github"}, "artifact_type": {"confidence": null, "enriched_at": "2026-08-29T17:20:12.522710+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "dcee470df8954629148ca2e53e70f8babfeeab97169745c1e00fb962dad3c777", "fetched_at": "2026-08-28T04:10:38.955510+00:00", "kind": "readme", "missing": false, "url": "https://github.com/OpenRLHF/OpenRLHF"}, {"content_hash": "d555b38ec2b2cb657f871fa64fa66c660d299caadc276af42634105dc26e9e31", "fetched_at": "2026-08-29T08:19:56.914305+00:00", "kind": "homepage", "missing": false, "url": "https://openrlhf.readthedocs.io/"}], "taxonomy_version": 1}, "description": {"kind": "observed", "observed_at": "2026-08-28T04:10:38.955510+00:00", "source": "github"}, "domain": {"confidence": null, "enriched_at": "2026-08-29T17:20:12.522710+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "dcee470df8954629148ca2e53e70f8babfeeab97169745c1e00fb962dad3c777", "fetched_at": "2026-08-28T04:10:38.955510+00:00", "kind": "readme", "missing": false, "url": "https://github.com/OpenRLHF/OpenRLHF"}, {"content_hash": "d555b38ec2b2cb657f871fa64fa66c660d299caadc276af42634105dc26e9e31", "fetched_at": "2026-08-29T08:19:56.914305+00:00", "kind": "homepage", "missing": false, "url": "https://openrlhf.readthedocs.io/"}], "taxonomy_version": 1}, "enriched": {"inputs": [], "kind": "computed", "method": "enrichment_status"}, "function": {"confidence": null, "enriched_at": "2026-08-29T17:20:12.522710+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "dcee470df8954629148ca2e53e70f8babfeeab97169745c1e00fb962dad3c777", "fetched_at": "2026-08-28T04:10:38.955510+00:00", "kind": "readme", "missing": false, "url": "https://github.com/OpenRLHF/OpenRLHF"}, {"content_hash": "d555b38ec2b2cb657f871fa64fa66c660d299caadc276af42634105dc26e9e31", "fetched_at": "2026-08-29T08:19:56.914305+00:00", "kind": "homepage", "missing": false, "url": "https://openrlhf.readthedocs.io/"}], "taxonomy_version": 1}, "health_score": {"inputs": ["days_since_push", "days_since_release", "archived"], "kind": "computed", "method": "health_v1"}, "homepage": {"kind": "observed", "observed_at": "2026-08-28T04:10:38.955510+00:00", "source": "github"}, "language": {"kind": "observed", "observed_at": "2026-08-28T04:10:38.955510+00:00", "source": "github"}, "license": {"kind": "observed", "observed_at": "2026-08-28T04:10:38.955510+00:00", "source": "github"}, "license_family": {"inputs": ["license"], "kind": "computed", "method": "license_family"}, "maturity": {"confidence": null, "enriched_at": "2026-08-29T17:20:12.522710+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "dcee470df8954629148ca2e53e70f8babfeeab97169745c1e00fb962dad3c777", "fetched_at": "2026-08-28T04:10:38.955510+00:00", "kind": "readme", "missing": false, "url": "https://github.com/OpenRLHF/OpenRLHF"}, {"content_hash": "d555b38ec2b2cb657f871fa64fa66c660d299caadc276af42634105dc26e9e31", "fetched_at": "2026-08-29T08:19:56.914305+00:00", "kind": "homepage", "missing": false, "url": "https://openrlhf.readthedocs.io/"}], "taxonomy_version": 1}, "member_repos": {"kind": "observed", "observed_at": "2026-08-28T04:10:38.955510+00:00", "source": "github"}, "name": {"kind": "observed", "observed_at": "2026-08-28T04:10:38.955510+00:00", "source": "github"}, "platform": {"confidence": null, "enriched_at": "2026-08-29T17:20:12.522710+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "dcee470df8954629148ca2e53e70f8babfeeab97169745c1e00fb962dad3c777", "fetched_at": "2026-08-28T04:10:38.955510+00:00", "kind": "readme", "missing": false, "url": "https://github.com/OpenRLHF/OpenRLHF"}, {"content_hash": "d555b38ec2b2cb657f871fa64fa66c660d299caadc276af42634105dc26e9e31", "fetched_at": "2026-08-29T08:19:56.914305+00:00", "kind": "homepage", "missing": false, "url": "https://openrlhf.readthedocs.io/"}], "taxonomy_version": 1}, "pushed_at": {"kind": "observed", "observed_at": "2026-08-28T04:10:38.955510+00:00", "source": "github"}, "repo": {"kind": "observed", "observed_at": "2026-08-28T04:10:38.955510+00:00", "source": "github"}, "stars": {"kind": "observed", "observed_at": "2026-08-28T04:10:38.955510+00:00", "source": "github"}, "tags": {"confidence": null, "enriched_at": "2026-08-29T17:20:12.522710+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "dcee470df8954629148ca2e53e70f8babfeeab97169745c1e00fb962dad3c777", "fetched_at": "2026-08-28T04:10:38.955510+00:00", "kind": "readme", "missing": false, "url": "https://github.com/OpenRLHF/OpenRLHF"}, {"content_hash": "d555b38ec2b2cb657f871fa64fa66c660d299caadc276af42634105dc26e9e31", "fetched_at": "2026-08-29T08:19:56.914305+00:00", "kind": "homepage", "missing": false, "url": "https://openrlhf.readthedocs.io/"}], "taxonomy_version": 1}, "topics": {"kind": "observed", "observed_at": "2026-08-28T04:10:38.955510+00:00", "source": "github"}, "urls": {"kind": "observed", "observed_at": "2026-08-28T04:10:38.955510+00:00", "source": "github"}, "use_cases": {"confidence": null, "enriched_at": "2026-08-29T17:20:12.522710+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "dcee470df8954629148ca2e53e70f8babfeeab97169745c1e00fb962dad3c777", "fetched_at": "2026-08-28T04:10:38.955510+00:00", "kind": "readme", "missing": false, "url": "https://github.com/OpenRLHF/OpenRLHF"}, {"content_hash": "d555b38ec2b2cb657f871fa64fa66c660d299caadc276af42634105dc26e9e31", "fetched_at": "2026-08-29T08:19:56.914305+00:00", "kind": "homepage", "missing": false, "url": "https://openrlhf.readthedocs.io/"}], "taxonomy_version": 1}, "what_it_is": {"confidence": null, "enriched_at": "2026-08-29T17:20:12.522710+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "dcee470df8954629148ca2e53e70f8babfeeab97169745c1e00fb962dad3c777", "fetched_at": "2026-08-28T04:10:38.955510+00:00", "kind": "readme", "missing": false, "url": "https://github.com/OpenRLHF/OpenRLHF"}, {"content_hash": "d555b38ec2b2cb657f871fa64fa66c660d299caadc276af42634105dc26e9e31", "fetched_at": "2026-08-29T08:19:56.914305+00:00", "kind": "homepage", "missing": false, "url": "https://openrlhf.readthedocs.io/"}], "taxonomy_version": 1}, "when_to_avoid": {"confidence": null, "enriched_at": "2026-08-29T17:20:12.522710+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "dcee470df8954629148ca2e53e70f8babfeeab97169745c1e00fb962dad3c777", "fetched_at": "2026-08-28T04:10:38.955510+00:00", "kind": "readme", "missing": false, "url": "https://github.com/OpenRLHF/OpenRLHF"}, {"content_hash": "d555b38ec2b2cb657f871fa64fa66c660d299caadc276af42634105dc26e9e31", "fetched_at": "2026-08-29T08:19:56.914305+00:00", "kind": "homepage", "missing": false, "url": "https://openrlhf.readthedocs.io/"}], "taxonomy_version": 1}, "when_to_choose": {"confidence": null, "enriched_at": "2026-08-29T17:20:12.522710+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "dcee470df8954629148ca2e53e70f8babfeeab97169745c1e00fb962dad3c777", "fetched_at": "2026-08-28T04:10:38.955510+00:00", "kind": "readme", "missing": false, "url": "https://github.com/OpenRLHF/OpenRLHF"}, {"content_hash": "d555b38ec2b2cb657f871fa64fa66c660d299caadc276af42634105dc26e9e31", "fetched_at": "2026-08-29T08:19:56.914305+00:00", "kind": "homepage", "missing": false, "url": "https://openrlhf.readthedocs.io/"}], "taxonomy_version": 1}}, "score": {"components": {"activity": 97, "longevity": 80, "rhythm": 85}, "computed_at": "2026-09-03T02:39:23.370411+00:00", "flags": [], "formula": "round(0.45*activity + 0.35*rhythm + 0.20*longevity); archived -> min(score, 10)", "inputs": {"age_days": 1131, "days_push": 20, "days_rel": 20, "gap_med": 5.0, "n_releases_24m": 73}, "score": 89, "version": 2}, "staleness": {"enrichment_outdated": false, "low_confidence": false, "scrape_days": 9, "stale_scrape": false}}