{"adoption": {"forks": 149, "observed_at": "2026-08-28T04:04:09.716827+00:00", "stars": 1259}, "canonical_url": "https://ross.abutalabs.com/products/following-instructions-human-feedback", "card": {"archived": true, "artifact_type": "learning-resource", "description": null, "domain": ["large-language-models", "machine-learning", "artificial-intelligence", "tutorials"], "enriched": true, "function": ["machine-learning", "llm-training", "reinforcement-learning", "nlp"], "health_score": 10, "homepage": null, "language": null, "license": null, "license_family": "other", "maturity": "maintenance", "member_repos": ["openai/following-instructions-human-feedback"], "name": "openai/following-instructions-human-feedback", "platform": ["python"], "pushed_at": "2022-12-11T19:58:53+00:00", "repo": "openai/following-instructions-human-feedback", "stars": 1259, "tags": ["rlhf", "instructgpt", "human-feedback", "model-alignment", "paper-companion", "model-card", "gpt-3", "fine-tuning", "natural-language-processing", "gpu", "linux"], "topics": [], "urls": [], "use_cases": ["understand how RLHF aligns language models with user intent", "read the InstructGPT model card and evaluation methodology", "study labeling instructions used for human feedback data collection", "compare GPT-3 and InstructGPT outputs on NLP benchmarks", "learn how supervised fine-tuning plus reward modeling reduces toxicity", "research alignment techniques for large language models"], "what_it_is": "The official companion repository for OpenAI's InstructGPT paper on aligning language models with human intent via reinforcement learning from human feedback (RLHF). It contains the model card, evaluation samples, and labeling instructions rather than runnable training code.", "when_to_avoid": ["you need runnable RLHF training code - this repo ships no implementation", "you want the actual InstructGPT weights or datasets - they are not included", "you need a maintained library - the repo is a static paper companion"], "when_to_choose": ["you want the authoritative artifacts and documentation behind the InstructGPT paper", "you are studying RLHF and model alignment from primary sources", "you need the official model card or labeling guidelines for citation or replication"]}, "data_as_of": "2026-08-30T08:39:29.467469+00:00", "members": [{"path": "/products/following-instructions-human-feedback", "repo": "openai/following-instructions-human-feedback", "role": "main", "score": 10}], "provenance": {"archived": {"kind": "observed", "observed_at": "2026-08-28T04:04:09.716827+00:00", "source": "github"}, "artifact_type": {"confidence": null, "enriched_at": "2026-08-30T05:05:31.638080+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "8ba860f7ffca02a147f4148589fc757391c94d6ba86fff94192de96c86b3ccb1", "fetched_at": "2026-08-28T04:04:09.716827+00:00", "kind": "readme", "missing": false, "url": "https://github.com/openai/following-instructions-human-feedback"}], "taxonomy_version": 1}, "description": {"kind": "observed", "observed_at": "2026-08-28T04:04:09.716827+00:00", "source": "github"}, "domain": {"confidence": null, "enriched_at": "2026-08-30T05:05:31.638080+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "8ba860f7ffca02a147f4148589fc757391c94d6ba86fff94192de96c86b3ccb1", "fetched_at": "2026-08-28T04:04:09.716827+00:00", "kind": "readme", "missing": false, "url": "https://github.com/openai/following-instructions-human-feedback"}], "taxonomy_version": 1}, "enriched": {"inputs": [], "kind": "computed", "method": "enrichment_status"}, "function": {"confidence": null, "enriched_at": "2026-08-30T05:05:31.638080+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "8ba860f7ffca02a147f4148589fc757391c94d6ba86fff94192de96c86b3ccb1", "fetched_at": "2026-08-28T04:04:09.716827+00:00", "kind": "readme", "missing": false, "url": "https://github.com/openai/following-instructions-human-feedback"}], "taxonomy_version": 1}, "health_score": {"inputs": ["days_since_push", "days_since_release", "archived"], "kind": "computed", "method": "health_v1"}, "homepage": {"kind": "observed", "observed_at": "2026-08-28T04:04:09.716827+00:00", "source": "github"}, "language": {"kind": "observed", "observed_at": "2026-08-28T04:04:09.716827+00:00", "source": "github"}, "license": {"kind": "observed", "observed_at": "2026-08-28T04:04:09.716827+00:00", "source": "github"}, "license_family": {"inputs": ["license"], "kind": "computed", "method": "license_family"}, "maturity": {"confidence": null, "enriched_at": "2026-08-30T05:05:31.638080+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "8ba860f7ffca02a147f4148589fc757391c94d6ba86fff94192de96c86b3ccb1", "fetched_at": "2026-08-28T04:04:09.716827+00:00", "kind": "readme", "missing": false, "url": "https://github.com/openai/following-instructions-human-feedback"}], "taxonomy_version": 1}, "member_repos": {"kind": "observed", "observed_at": "2026-08-28T04:04:09.716827+00:00", "source": "github"}, "name": {"kind": "observed", "observed_at": "2026-08-28T04:04:09.716827+00:00", "source": "github"}, "platform": {"confidence": null, "enriched_at": "2026-08-30T05:05:31.638080+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "8ba860f7ffca02a147f4148589fc757391c94d6ba86fff94192de96c86b3ccb1", "fetched_at": "2026-08-28T04:04:09.716827+00:00", "kind": "readme", "missing": false, "url": "https://github.com/openai/following-instructions-human-feedback"}], "taxonomy_version": 1}, "pushed_at": {"kind": "observed", "observed_at": "2026-08-28T04:04:09.716827+00:00", "source": "github"}, "repo": {"kind": "observed", "observed_at": "2026-08-28T04:04:09.716827+00:00", "source": "github"}, "stars": {"kind": "observed", "observed_at": "2026-08-28T04:04:09.716827+00:00", "source": "github"}, "tags": {"confidence": null, "enriched_at": "2026-08-30T05:05:31.638080+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "8ba860f7ffca02a147f4148589fc757391c94d6ba86fff94192de96c86b3ccb1", "fetched_at": "2026-08-28T04:04:09.716827+00:00", "kind": "readme", "missing": false, "url": "https://github.com/openai/following-instructions-human-feedback"}], "taxonomy_version": 1}, "topics": {"kind": "observed", "observed_at": "2026-08-28T04:04:09.716827+00:00", "source": "github"}, "urls": {"kind": "observed", "observed_at": "2026-08-28T04:04:09.716827+00:00", "source": "github"}, "use_cases": {"confidence": null, "enriched_at": "2026-08-30T05:05:31.638080+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "8ba860f7ffca02a147f4148589fc757391c94d6ba86fff94192de96c86b3ccb1", "fetched_at": "2026-08-28T04:04:09.716827+00:00", "kind": "readme", "missing": false, "url": "https://github.com/openai/following-instructions-human-feedback"}], "taxonomy_version": 1}, "what_it_is": {"confidence": null, "enriched_at": "2026-08-30T05:05:31.638080+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "8ba860f7ffca02a147f4148589fc757391c94d6ba86fff94192de96c86b3ccb1", "fetched_at": "2026-08-28T04:04:09.716827+00:00", "kind": "readme", "missing": false, "url": "https://github.com/openai/following-instructions-human-feedback"}], "taxonomy_version": 1}, "when_to_avoid": {"confidence": null, "enriched_at": "2026-08-30T05:05:31.638080+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "8ba860f7ffca02a147f4148589fc757391c94d6ba86fff94192de96c86b3ccb1", "fetched_at": "2026-08-28T04:04:09.716827+00:00", "kind": "readme", "missing": false, "url": "https://github.com/openai/following-instructions-human-feedback"}], "taxonomy_version": 1}, "when_to_choose": {"confidence": null, "enriched_at": "2026-08-30T05:05:31.638080+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "8ba860f7ffca02a147f4148589fc757391c94d6ba86fff94192de96c86b3ccb1", "fetched_at": "2026-08-28T04:04:09.716827+00:00", "kind": "readme", "missing": false, "url": "https://github.com/openai/following-instructions-human-feedback"}], "taxonomy_version": 1}}, "score": {"components": {"activity": 0, "longevity": 100, "rhythm": 35}, "computed_at": "2026-09-03T02:20:16.233290+00:00", "flags": ["no_releases", "archived", "no_license"], "formula": "round(0.45*activity + 0.35*rhythm + 0.20*longevity); archived -> min(score, 10)", "inputs": {"age_days": 1682, "days_push": 1361, "days_rel": null, "gap_med": null, "n_releases_24m": 0}, "score": 10, "version": 2}, "staleness": {"enrichment_outdated": false, "low_confidence": false, "scrape_days": 9, "stale_scrape": false}}