{"adoption": {"forks": 370, "observed_at": "2026-08-28T04:06:23.722555+00:00", "stars": 2189}, "canonical_url": "https://ross.abutalabs.com/products/spark-vllm-docker", "card": {"archived": false, "artifact_type": "infra-config", "description": "Docker configuration for running VLLM on dual DGX Sparks", "domain": ["large-language-models", "gpu-computing", "infrastructure-as-code", "self-hosted"], "enriched": true, "function": ["llm-inference", "container-runtime", "deployment", "gpu-computing", "developer-tools"], "health_score": 100, "homepage": null, "language": "Shell", "license": "MIT", "license_family": "permissive", "maturity": "active", "member_repos": ["eugr/spark-vllm-docker"], "name": "eugr/spark-vllm-docker", "platform": ["self-hosted"], "pushed_at": "2026-08-26T12:12:24+00:00", "repo": "eugr/spark-vllm-docker", "stars": 2189, "tags": ["vllm", "dgx-spark", "multi-node", "infiniband", "nccl", "ray", "docker-compose", "nvidia", "containers", "docker", "linux", "gpu"], "topics": [], "urls": [], "use_cases": ["run vllm on dgx spark", "set up multi-node llm inference cluster", "serve large language models across two dgx sparks", "dockerize vllm with infiniband support", "deploy vllm with ray distributed backend", "benchmark llm inference on spark hardware", "download and load models fast with fastsafetensors"], "what_it_is": "A Docker configuration and set of shell scripts for running vLLM inference on NVIDIA DGX Spark hardware, from single nodes to multi-node clusters using Ray or PyTorch distributed. It includes support for InfiniBand/RDMA networking, fast model loading, and tested nightly Docker images.", "when_to_avoid": ["you run inference on non-DGX-Spark GPUs or cloud instances", "you need a general-purpose vLLM deployment on Kubernetes", "you want a managed inference service rather than self-hosted infrastructure", "you need training or fine-tuning rather than inference"], "when_to_choose": ["you own one or more DGX Spark machines and want optimized vLLM inference", "you need multi-node LLM serving with InfiniBand/RDMA or QSFP networking", "you want pre-tested nightly Docker images instead of hand-building vLLM", "you need cluster orchestration via Ray or native PyTorch distributed mode"]}, "data_as_of": "2026-08-30T08:39:29.467469+00:00", "members": [{"path": "/products/spark-vllm-docker", "repo": "eugr/spark-vllm-docker", "role": "main", "score": 83}], "provenance": {"archived": {"kind": "observed", "observed_at": "2026-08-28T04:06:23.722555+00:00", "source": "github"}, "artifact_type": {"confidence": null, "enriched_at": "2026-08-30T02:47:43.752720+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "fc31885bb6a99856a8fc68ffba356a273005573df8576176fdba8852c8172842", "fetched_at": "2026-08-28T04:06:23.722555+00:00", "kind": "readme", "missing": false, "url": "https://github.com/eugr/spark-vllm-docker"}], "taxonomy_version": 1}, "description": {"kind": "observed", "observed_at": "2026-08-28T04:06:23.722555+00:00", "source": "github"}, "domain": {"confidence": null, "enriched_at": "2026-08-30T02:47:43.752720+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "fc31885bb6a99856a8fc68ffba356a273005573df8576176fdba8852c8172842", "fetched_at": "2026-08-28T04:06:23.722555+00:00", "kind": "readme", "missing": false, "url": "https://github.com/eugr/spark-vllm-docker"}], "taxonomy_version": 1}, "enriched": {"inputs": [], "kind": "computed", "method": "enrichment_status"}, "function": {"confidence": null, "enriched_at": "2026-08-30T02:47:43.752720+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "fc31885bb6a99856a8fc68ffba356a273005573df8576176fdba8852c8172842", "fetched_at": "2026-08-28T04:06:23.722555+00:00", "kind": "readme", "missing": false, "url": "https://github.com/eugr/spark-vllm-docker"}], "taxonomy_version": 1}, "health_score": {"inputs": ["days_since_push", "days_since_release", "archived"], "kind": "computed", "method": "health_v1"}, "homepage": {"kind": "observed", "observed_at": "2026-08-28T04:06:23.722555+00:00", "source": "github"}, "language": {"kind": "observed", "observed_at": "2026-08-28T04:06:23.722555+00:00", "source": "github"}, "license": {"kind": "observed", "observed_at": "2026-08-28T04:06:23.722555+00:00", "source": "github"}, "license_family": {"inputs": ["license"], "kind": "computed", "method": "license_family"}, "maturity": {"confidence": null, "enriched_at": "2026-08-30T02:47:43.752720+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "fc31885bb6a99856a8fc68ffba356a273005573df8576176fdba8852c8172842", "fetched_at": "2026-08-28T04:06:23.722555+00:00", "kind": "readme", "missing": false, "url": "https://github.com/eugr/spark-vllm-docker"}], "taxonomy_version": 1}, "member_repos": {"kind": "observed", "observed_at": "2026-08-28T04:06:23.722555+00:00", "source": "github"}, "name": {"kind": "observed", "observed_at": "2026-08-28T04:06:23.722555+00:00", "source": "github"}, "platform": {"confidence": null, "enriched_at": "2026-08-30T02:47:43.752720+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "fc31885bb6a99856a8fc68ffba356a273005573df8576176fdba8852c8172842", "fetched_at": "2026-08-28T04:06:23.722555+00:00", "kind": "readme", "missing": false, "url": "https://github.com/eugr/spark-vllm-docker"}], "taxonomy_version": 1}, "pushed_at": {"kind": "observed", "observed_at": "2026-08-28T04:06:23.722555+00:00", "source": "github"}, "repo": {"kind": "observed", "observed_at": "2026-08-28T04:06:23.722555+00:00", "source": "github"}, "stars": {"kind": "observed", "observed_at": "2026-08-28T04:06:23.722555+00:00", "source": "github"}, "tags": {"confidence": null, "enriched_at": "2026-08-30T02:47:43.752720+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "fc31885bb6a99856a8fc68ffba356a273005573df8576176fdba8852c8172842", "fetched_at": "2026-08-28T04:06:23.722555+00:00", "kind": "readme", "missing": false, "url": "https://github.com/eugr/spark-vllm-docker"}], "taxonomy_version": 1}, "topics": {"kind": "observed", "observed_at": "2026-08-28T04:06:23.722555+00:00", "source": "github"}, "urls": {"kind": "observed", "observed_at": "2026-08-28T04:06:23.722555+00:00", "source": "github"}, "use_cases": {"confidence": null, "enriched_at": "2026-08-30T02:47:43.752720+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "fc31885bb6a99856a8fc68ffba356a273005573df8576176fdba8852c8172842", "fetched_at": "2026-08-28T04:06:23.722555+00:00", "kind": "readme", "missing": false, "url": "https://github.com/eugr/spark-vllm-docker"}], "taxonomy_version": 1}, "what_it_is": {"confidence": null, "enriched_at": "2026-08-30T02:47:43.752720+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "fc31885bb6a99856a8fc68ffba356a273005573df8576176fdba8852c8172842", "fetched_at": "2026-08-28T04:06:23.722555+00:00", "kind": "readme", "missing": false, "url": "https://github.com/eugr/spark-vllm-docker"}], "taxonomy_version": 1}, "when_to_avoid": {"confidence": null, "enriched_at": "2026-08-30T02:47:43.752720+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "fc31885bb6a99856a8fc68ffba356a273005573df8576176fdba8852c8172842", "fetched_at": "2026-08-28T04:06:23.722555+00:00", "kind": "readme", "missing": false, "url": "https://github.com/eugr/spark-vllm-docker"}], "taxonomy_version": 1}, "when_to_choose": {"confidence": null, "enriched_at": "2026-08-30T02:47:43.752720+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "fc31885bb6a99856a8fc68ffba356a273005573df8576176fdba8852c8172842", "fetched_at": "2026-08-28T04:06:23.722555+00:00", "kind": "readme", "missing": false, "url": "https://github.com/eugr/spark-vllm-docker"}], "taxonomy_version": 1}}, "score": {"components": {"activity": 99, "longevity": 20, "rhythm": 99}, "computed_at": "2026-09-03T02:20:16.233290+00:00", "flags": [], "formula": "round(0.45*activity + 0.35*rhythm + 0.20*longevity); archived -> min(score, 10)", "inputs": {"age_days": 281, "days_push": 7, "days_rel": 7, "gap_med": 1, "n_releases_24m": 2}, "score": 83, "version": 2}, "staleness": {"enrichment_outdated": false, "low_confidence": false, "scrape_days": 9, "stale_scrape": false}}