{"adoption": {"forks": 1554, "observed_at": "2026-08-28T04:09:42.825456+00:00", "stars": 6369}, "canonical_url": "https://ross.abutalabs.com/products/vllm-omni", "card": {"archived": false, "artifact_type": "framework", "description": "A framework for efficient model inference with omni-modality models", "domain": ["large-language-models", "machine-learning", "artificial-intelligence", "image-processing", "speech-processing", "robotics", "gpu-computing"], "enriched": true, "function": ["llm-inference", "machine-learning", "image-processing", "audio-processing", "video-processing", "tts", "speech-recognition", "http-server", "api-framework"], "health_score": 100, "homepage": "https://docs.vllm.ai/projects/vllm-omni", "language": "Python", "license": "Apache-2.0", "license_family": "permissive", "maturity": "active", "member_repos": ["vllm-project/vllm-omni"], "name": "vllm-project/vllm-omni", "platform": ["python", "cloud"], "pushed_at": "2026-08-26T20:19:21+00:00", "repo": "vllm-project/vllm-omni", "stars": 6369, "tags": ["diffusion", "multimodal", "model-serving", "openai-compatible-api", "dit", "world-model", "robot-policy", "vllm", "audio", "video", "linux", "gpu", "docker"], "topics": ["diffusion", "inference", "model-serving", "pytorch", "transformer", "audio-generation", "image-generation", "multimodal", "video-generation", "world-model"], "urls": [], "use_cases": ["serve a text-to-image diffusion model behind an OpenAI-compatible API", "run offline batched inference with Qwen3-Omni or MiniCPM-o", "deploy TTS models like CosyVoice3 with streaming audio output", "serve video generation models like Wan2.2 or MiniMax H3", "host robot-policy and action models for robotics inference", "run full-duplex realtime voice serving with streaming audio input and output"], "what_it_is": "vLLM-Omni is a Python framework extending vLLM for efficient inference and serving of omni-modality models, including diffusion transformers, TTS, image/video generation, and robot-policy models. It provides pipelined stage execution, distributed parallelism, streaming outputs, and an OpenAI-compatible API server.", "when_to_avoid": ["you only need plain text LLM inference, where core vLLM suffices", "you need Windows or macOS support, since it targets Linux", "you want a lightweight single-model pipeline without distributed serving complexity"], "when_to_choose": ["you need high-throughput serving of multimodal, diffusion, TTS, or action models on GPU", "you want vLLM-style performance (KV cache, batching, parallelism) for non-autoregressive models", "you need an OpenAI-compatible server with streaming multimodal outputs"]}, "data_as_of": "2026-08-30T08:39:29.467469+00:00", "members": [{"path": "/products/vllm-omni", "repo": "vllm-project/vllm-omni", "role": "main", "score": 83}], "provenance": {"archived": {"kind": "observed", "observed_at": "2026-08-28T04:09:42.825456+00:00", "source": "github"}, "artifact_type": {"confidence": null, "enriched_at": "2026-08-29T17:45:12.878419+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "21fd672b01650d94e91e706a09d18ebd41caea83ee218c25b841d4f88fa5a636", "fetched_at": "2026-08-28T04:09:42.825456+00:00", "kind": "readme", "missing": false, "url": "https://github.com/vllm-project/vllm-omni"}, {"content_hash": "50487ae1953f20ea1e8e6efe640497e15779d0370ec886c205cbdcf731b3cc4d", "fetched_at": "2026-08-29T08:42:13.857704+00:00", "kind": "homepage", "missing": false, "url": "https://docs.vllm.ai/projects/vllm-omni"}, {"content_hash": "46d1086ad5156a590a8f54a94411092b705fd40f9b183a396020301815a6234e", "fetched_at": "2026-08-29T08:42:13.866648+00:00", "kind": "site_page", "missing": false, "url": "https://docs.vllm.ai/projects/vllm-omni/en/latest/getting_started/quickstart"}, {"content_hash": "8e7ecb34ac228f3fb8545447125926bfd705728d66ef9d6fab9fa09047d0ee42", "fetched_at": "2026-08-29T08:42:13.868730+00:00", "kind": "registry_pypi", "missing": false, "url": "https://pypi.org/pypi/vllm-omni/json"}], "taxonomy_version": 1}, "description": {"kind": "observed", "observed_at": "2026-08-28T04:09:42.825456+00:00", "source": "github"}, "domain": {"confidence": null, "enriched_at": "2026-08-29T17:45:12.878419+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "21fd672b01650d94e91e706a09d18ebd41caea83ee218c25b841d4f88fa5a636", "fetched_at": "2026-08-28T04:09:42.825456+00:00", "kind": "readme", "missing": false, "url": "https://github.com/vllm-project/vllm-omni"}, {"content_hash": "50487ae1953f20ea1e8e6efe640497e15779d0370ec886c205cbdcf731b3cc4d", "fetched_at": "2026-08-29T08:42:13.857704+00:00", "kind": "homepage", "missing": false, "url": "https://docs.vllm.ai/projects/vllm-omni"}, {"content_hash": "46d1086ad5156a590a8f54a94411092b705fd40f9b183a396020301815a6234e", "fetched_at": "2026-08-29T08:42:13.866648+00:00", "kind": "site_page", "missing": false, "url": "https://docs.vllm.ai/projects/vllm-omni/en/latest/getting_started/quickstart"}, {"content_hash": "8e7ecb34ac228f3fb8545447125926bfd705728d66ef9d6fab9fa09047d0ee42", "fetched_at": "2026-08-29T08:42:13.868730+00:00", "kind": "registry_pypi", "missing": false, "url": "https://pypi.org/pypi/vllm-omni/json"}], "taxonomy_version": 1}, "enriched": {"inputs": [], "kind": "computed", "method": "enrichment_status"}, "function": {"confidence": null, "enriched_at": "2026-08-29T17:45:12.878419+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "21fd672b01650d94e91e706a09d18ebd41caea83ee218c25b841d4f88fa5a636", "fetched_at": "2026-08-28T04:09:42.825456+00:00", "kind": "readme", "missing": false, "url": "https://github.com/vllm-project/vllm-omni"}, {"content_hash": "50487ae1953f20ea1e8e6efe640497e15779d0370ec886c205cbdcf731b3cc4d", "fetched_at": "2026-08-29T08:42:13.857704+00:00", "kind": "homepage", "missing": false, "url": "https://docs.vllm.ai/projects/vllm-omni"}, {"content_hash": "46d1086ad5156a590a8f54a94411092b705fd40f9b183a396020301815a6234e", "fetched_at": "2026-08-29T08:42:13.866648+00:00", "kind": "site_page", "missing": false, "url": "https://docs.vllm.ai/projects/vllm-omni/en/latest/getting_started/quickstart"}, {"content_hash": "8e7ecb34ac228f3fb8545447125926bfd705728d66ef9d6fab9fa09047d0ee42", "fetched_at": "2026-08-29T08:42:13.868730+00:00", "kind": "registry_pypi", "missing": false, "url": "https://pypi.org/pypi/vllm-omni/json"}], "taxonomy_version": 1}, "health_score": {"inputs": ["days_since_push", "days_since_release", "archived"], "kind": "computed", "method": "health_v1"}, "homepage": {"kind": "observed", "observed_at": "2026-08-28T04:09:42.825456+00:00", "source": "github"}, "language": {"kind": "observed", "observed_at": "2026-08-28T04:09:42.825456+00:00", "source": "github"}, "license": {"kind": "observed", "observed_at": "2026-08-28T04:09:42.825456+00:00", "source": "github"}, "license_family": {"inputs": ["license"], "kind": "computed", "method": "license_family"}, "maturity": {"confidence": null, "enriched_at": "2026-08-29T17:45:12.878419+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "21fd672b01650d94e91e706a09d18ebd41caea83ee218c25b841d4f88fa5a636", "fetched_at": "2026-08-28T04:09:42.825456+00:00", "kind": "readme", "missing": false, "url": "https://github.com/vllm-project/vllm-omni"}, {"content_hash": "50487ae1953f20ea1e8e6efe640497e15779d0370ec886c205cbdcf731b3cc4d", "fetched_at": "2026-08-29T08:42:13.857704+00:00", "kind": "homepage", "missing": false, "url": "https://docs.vllm.ai/projects/vllm-omni"}, {"content_hash": "46d1086ad5156a590a8f54a94411092b705fd40f9b183a396020301815a6234e", "fetched_at": "2026-08-29T08:42:13.866648+00:00", "kind": "site_page", "missing": false, "url": "https://docs.vllm.ai/projects/vllm-omni/en/latest/getting_started/quickstart"}, {"content_hash": "8e7ecb34ac228f3fb8545447125926bfd705728d66ef9d6fab9fa09047d0ee42", "fetched_at": "2026-08-29T08:42:13.868730+00:00", "kind": "registry_pypi", "missing": false, "url": "https://pypi.org/pypi/vllm-omni/json"}], "taxonomy_version": 1}, "member_repos": {"kind": "observed", "observed_at": "2026-08-28T04:09:42.825456+00:00", "source": "github"}, "name": {"kind": "observed", "observed_at": "2026-08-28T04:09:42.825456+00:00", "source": "github"}, "platform": {"confidence": null, "enriched_at": "2026-08-29T17:45:12.878419+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "21fd672b01650d94e91e706a09d18ebd41caea83ee218c25b841d4f88fa5a636", "fetched_at": "2026-08-28T04:09:42.825456+00:00", "kind": "readme", "missing": false, "url": "https://github.com/vllm-project/vllm-omni"}, {"content_hash": "50487ae1953f20ea1e8e6efe640497e15779d0370ec886c205cbdcf731b3cc4d", "fetched_at": "2026-08-29T08:42:13.857704+00:00", "kind": "homepage", "missing": false, "url": "https://docs.vllm.ai/projects/vllm-omni"}, {"content_hash": "46d1086ad5156a590a8f54a94411092b705fd40f9b183a396020301815a6234e", "fetched_at": "2026-08-29T08:42:13.866648+00:00", "kind": "site_page", "missing": false, "url": "https://docs.vllm.ai/projects/vllm-omni/en/latest/getting_started/quickstart"}, {"content_hash": "8e7ecb34ac228f3fb8545447125926bfd705728d66ef9d6fab9fa09047d0ee42", "fetched_at": "2026-08-29T08:42:13.868730+00:00", "kind": "registry_pypi", "missing": false, "url": "https://pypi.org/pypi/vllm-omni/json"}], "taxonomy_version": 1}, "pushed_at": {"kind": "observed", "observed_at": "2026-08-28T04:09:42.825456+00:00", "source": "github"}, "repo": {"kind": "observed", "observed_at": "2026-08-28T04:09:42.825456+00:00", "source": "github"}, "stars": {"kind": "observed", "observed_at": "2026-08-28T04:09:42.825456+00:00", "source": "github"}, "tags": {"confidence": null, "enriched_at": "2026-08-29T17:45:12.878419+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "21fd672b01650d94e91e706a09d18ebd41caea83ee218c25b841d4f88fa5a636", "fetched_at": "2026-08-28T04:09:42.825456+00:00", "kind": "readme", "missing": false, "url": "https://github.com/vllm-project/vllm-omni"}, {"content_hash": "50487ae1953f20ea1e8e6efe640497e15779d0370ec886c205cbdcf731b3cc4d", "fetched_at": "2026-08-29T08:42:13.857704+00:00", "kind": "homepage", "missing": false, "url": "https://docs.vllm.ai/projects/vllm-omni"}, {"content_hash": "46d1086ad5156a590a8f54a94411092b705fd40f9b183a396020301815a6234e", "fetched_at": "2026-08-29T08:42:13.866648+00:00", "kind": "site_page", "missing": false, "url": "https://docs.vllm.ai/projects/vllm-omni/en/latest/getting_started/quickstart"}, {"content_hash": "8e7ecb34ac228f3fb8545447125926bfd705728d66ef9d6fab9fa09047d0ee42", "fetched_at": "2026-08-29T08:42:13.868730+00:00", "kind": "registry_pypi", "missing": false, "url": "https://pypi.org/pypi/vllm-omni/json"}], "taxonomy_version": 1}, "topics": {"kind": "observed", "observed_at": "2026-08-28T04:09:42.825456+00:00", "source": "github"}, "urls": {"kind": "observed", "observed_at": "2026-08-28T04:09:42.825456+00:00", "source": "github"}, "use_cases": {"confidence": null, "enriched_at": "2026-08-29T17:45:12.878419+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "21fd672b01650d94e91e706a09d18ebd41caea83ee218c25b841d4f88fa5a636", "fetched_at": "2026-08-28T04:09:42.825456+00:00", "kind": "readme", "missing": false, "url": "https://github.com/vllm-project/vllm-omni"}, {"content_hash": "50487ae1953f20ea1e8e6efe640497e15779d0370ec886c205cbdcf731b3cc4d", "fetched_at": "2026-08-29T08:42:13.857704+00:00", "kind": "homepage", "missing": false, "url": "https://docs.vllm.ai/projects/vllm-omni"}, {"content_hash": "46d1086ad5156a590a8f54a94411092b705fd40f9b183a396020301815a6234e", "fetched_at": "2026-08-29T08:42:13.866648+00:00", "kind": "site_page", "missing": false, "url": "https://docs.vllm.ai/projects/vllm-omni/en/latest/getting_started/quickstart"}, {"content_hash": "8e7ecb34ac228f3fb8545447125926bfd705728d66ef9d6fab9fa09047d0ee42", "fetched_at": "2026-08-29T08:42:13.868730+00:00", "kind": "registry_pypi", "missing": false, "url": "https://pypi.org/pypi/vllm-omni/json"}], "taxonomy_version": 1}, "what_it_is": {"confidence": null, "enriched_at": "2026-08-29T17:45:12.878419+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "21fd672b01650d94e91e706a09d18ebd41caea83ee218c25b841d4f88fa5a636", "fetched_at": "2026-08-28T04:09:42.825456+00:00", "kind": "readme", "missing": false, "url": "https://github.com/vllm-project/vllm-omni"}, {"content_hash": "50487ae1953f20ea1e8e6efe640497e15779d0370ec886c205cbdcf731b3cc4d", "fetched_at": "2026-08-29T08:42:13.857704+00:00", "kind": "homepage", "missing": false, "url": "https://docs.vllm.ai/projects/vllm-omni"}, {"content_hash": "46d1086ad5156a590a8f54a94411092b705fd40f9b183a396020301815a6234e", "fetched_at": "2026-08-29T08:42:13.866648+00:00", "kind": "site_page", "missing": false, "url": "https://docs.vllm.ai/projects/vllm-omni/en/latest/getting_started/quickstart"}, {"content_hash": "8e7ecb34ac228f3fb8545447125926bfd705728d66ef9d6fab9fa09047d0ee42", "fetched_at": "2026-08-29T08:42:13.868730+00:00", "kind": "registry_pypi", "missing": false, "url": "https://pypi.org/pypi/vllm-omni/json"}], "taxonomy_version": 1}, "when_to_avoid": {"confidence": null, "enriched_at": "2026-08-29T17:45:12.878419+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "21fd672b01650d94e91e706a09d18ebd41caea83ee218c25b841d4f88fa5a636", "fetched_at": "2026-08-28T04:09:42.825456+00:00", "kind": "readme", "missing": false, "url": "https://github.com/vllm-project/vllm-omni"}, {"content_hash": "50487ae1953f20ea1e8e6efe640497e15779d0370ec886c205cbdcf731b3cc4d", "fetched_at": "2026-08-29T08:42:13.857704+00:00", "kind": "homepage", "missing": false, "url": "https://docs.vllm.ai/projects/vllm-omni"}, {"content_hash": "46d1086ad5156a590a8f54a94411092b705fd40f9b183a396020301815a6234e", "fetched_at": "2026-08-29T08:42:13.866648+00:00", "kind": "site_page", "missing": false, "url": "https://docs.vllm.ai/projects/vllm-omni/en/latest/getting_started/quickstart"}, {"content_hash": "8e7ecb34ac228f3fb8545447125926bfd705728d66ef9d6fab9fa09047d0ee42", "fetched_at": "2026-08-29T08:42:13.868730+00:00", "kind": "registry_pypi", "missing": false, "url": "https://pypi.org/pypi/vllm-omni/json"}], "taxonomy_version": 1}, "when_to_choose": {"confidence": null, "enriched_at": "2026-08-29T17:45:12.878419+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "21fd672b01650d94e91e706a09d18ebd41caea83ee218c25b841d4f88fa5a636", "fetched_at": "2026-08-28T04:09:42.825456+00:00", "kind": "readme", "missing": false, "url": "https://github.com/vllm-project/vllm-omni"}, {"content_hash": "50487ae1953f20ea1e8e6efe640497e15779d0370ec886c205cbdcf731b3cc4d", "fetched_at": "2026-08-29T08:42:13.857704+00:00", "kind": "homepage", "missing": false, "url": "https://docs.vllm.ai/projects/vllm-omni"}, {"content_hash": "46d1086ad5156a590a8f54a94411092b705fd40f9b183a396020301815a6234e", "fetched_at": "2026-08-29T08:42:13.866648+00:00", "kind": "site_page", "missing": false, "url": "https://docs.vllm.ai/projects/vllm-omni/en/latest/getting_started/quickstart"}, {"content_hash": "8e7ecb34ac228f3fb8545447125926bfd705728d66ef9d6fab9fa09047d0ee42", "fetched_at": "2026-08-29T08:42:13.868730+00:00", "kind": "registry_pypi", "missing": false, "url": "https://pypi.org/pypi/vllm-omni/json"}], "taxonomy_version": 1}}, "score": {"components": {"activity": 99, "longevity": 25, "rhythm": 96}, "computed_at": "2026-09-03T02:20:16.233290+00:00", "flags": [], "formula": "round(0.45*activity + 0.35*rhythm + 0.20*longevity); archived -> min(score, 10)", "inputs": {"age_days": 357, "days_push": 7, "days_rel": 30, "gap_med": 28, "n_releases_24m": 8}, "score": 83, "version": 2}, "staleness": {"enrichment_outdated": false, "low_confidence": false, "scrape_days": 9, "stale_scrape": false}}