{"adoption": {"forks": 812, "observed_at": "2026-08-28T04:10:28.607821+00:00", "stars": 9151}, "canonical_url": "https://ross.abutalabs.com/products/sensevoice", "card": {"archived": false, "artifact_type": "library", "description": "Open-source SenseVoiceSmall model for Mandarin, Cantonese, English, Japanese, and Korean ASR, language ID, emotion recognition, and audio event detection.", "domain": ["speech-processing", "machine-learning"], "enriched": true, "function": ["speech-recognition", "audio-processing", "machine-learning", "llm-inference"], "health_score": 99, "homepage": "https://huggingface.co/spaces/FunAudioLLM/SenseVoice", "language": "C", "license": "MIT", "license_family": "permissive", "maturity": "active", "member_repos": ["QwenAudio/SenseVoice"], "name": "QwenAudio/SenseVoice", "platform": ["python", "cpp", "cross-platform"], "pushed_at": "2026-08-18T02:13:34+00:00", "repo": "QwenAudio/SenseVoice", "stars": 9151, "tags": ["asr", "speech-to-text", "emotion-recognition", "audio-event-detection", "language-identification", "multilingual", "funasr", "whisper-alternative", "transcription", "cantonese", "audio", "natural-language-processing", "gpu"], "topics": ["asr", "speech-recognition", "speech-to-text", "cross-lingual", "pytorch", "speech-emotion-recognition", "multilingual", "audio-analysis", "audio-event-detection", "emotion-detection", "voice-ai", "whisper-alternative", "funasr", "language-identification", "sensevoice", "transcription", "cantonese", "llama-cpp", "multilingual-asr", "speech-understanding"], "urls": [], "use_cases": ["transcribe speech to text in Mandarin, Cantonese, English, Japanese, or Korean", "detect emotion from speech audio", "identify the spoken language of an audio clip", "detect audio events like applause, laughter, or coughing", "find a fast low-latency alternative to Whisper for speech recognition", "finetune an ASR model on domain-specific audio samples"], "what_it_is": "SenseVoice is an open-source speech foundation model (SenseVoiceSmall) providing multilingual ASR, spoken language identification, speech emotion recognition, and audio event detection for Mandarin, Cantonese, English, Japanese, and Korean. It uses a non-autoregressive end-to-end framework for low-latency inference and integrates with FunASR for deployment and finetuning.", "when_to_avoid": ["you need speaker diarization as a single-model output (it requires composing separate FunASR VAD and CAM++ pipelines)", "you need ASR for languages beyond the five supported by the released checkpoint", "you need a tiny embedded deployment without GPU or C++/llama.cpp tooling"], "when_to_choose": ["you need fast, accurate multilingual ASR with extra emotion and audio-event tags", "your target languages are Mandarin, Cantonese, English, Japanese, or Korean", "you want low-latency non-autoregressive inference or easy finetuning via FunASR"]}, "data_as_of": "2026-08-30T08:39:29.467469+00:00", "members": [{"path": "/products/sensevoice", "repo": "QwenAudio/SenseVoice", "role": "main", "score": 88}], "provenance": {"archived": {"kind": "observed", "observed_at": "2026-08-28T04:10:28.607821+00:00", "source": "github"}, "artifact_type": {"confidence": null, "enriched_at": "2026-08-29T17:23:24.024621+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "86e30a2f1672c40137c6d7e40a40ac09ce3996e7c06873d430c40dbec21f8e4b", "fetched_at": "2026-08-28T04:10:28.607821+00:00", "kind": "readme", "missing": false, "url": "https://github.com/QwenAudio/SenseVoice"}, {"content_hash": "00799e0fcae6da1dc6bb276b1e53a07b5b1fd10d45d41f64f0bc68e66aa6d295", "fetched_at": "2026-08-29T08:23:27.860255+00:00", "kind": "homepage", "missing": false, "url": "https://huggingface.co/spaces/FunAudioLLM/SenseVoice"}], "taxonomy_version": 1}, "description": {"kind": "observed", "observed_at": "2026-08-28T04:10:28.607821+00:00", "source": "github"}, "domain": {"confidence": null, "enriched_at": "2026-08-29T17:23:24.024621+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "86e30a2f1672c40137c6d7e40a40ac09ce3996e7c06873d430c40dbec21f8e4b", "fetched_at": "2026-08-28T04:10:28.607821+00:00", "kind": "readme", "missing": false, "url": "https://github.com/QwenAudio/SenseVoice"}, {"content_hash": "00799e0fcae6da1dc6bb276b1e53a07b5b1fd10d45d41f64f0bc68e66aa6d295", "fetched_at": "2026-08-29T08:23:27.860255+00:00", "kind": "homepage", "missing": false, "url": "https://huggingface.co/spaces/FunAudioLLM/SenseVoice"}], "taxonomy_version": 1}, "enriched": {"inputs": [], "kind": "computed", "method": "enrichment_status"}, "function": {"confidence": null, "enriched_at": "2026-08-29T17:23:24.024621+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "86e30a2f1672c40137c6d7e40a40ac09ce3996e7c06873d430c40dbec21f8e4b", "fetched_at": "2026-08-28T04:10:28.607821+00:00", "kind": "readme", "missing": false, "url": "https://github.com/QwenAudio/SenseVoice"}, {"content_hash": "00799e0fcae6da1dc6bb276b1e53a07b5b1fd10d45d41f64f0bc68e66aa6d295", "fetched_at": "2026-08-29T08:23:27.860255+00:00", "kind": "homepage", "missing": false, "url": "https://huggingface.co/spaces/FunAudioLLM/SenseVoice"}], "taxonomy_version": 1}, "health_score": {"inputs": ["days_since_push", "days_since_release", "archived"], "kind": "computed", "method": "health_v1"}, "homepage": {"kind": "observed", "observed_at": "2026-08-28T04:10:28.607821+00:00", "source": "github"}, "language": {"kind": "observed", "observed_at": "2026-08-28T04:10:28.607821+00:00", "source": "github"}, "license": {"kind": "observed", "observed_at": "2026-08-28T04:10:28.607821+00:00", "source": "github"}, "license_family": {"inputs": ["license"], "kind": "computed", "method": "license_family"}, "maturity": {"confidence": null, "enriched_at": "2026-08-29T17:23:24.024621+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "86e30a2f1672c40137c6d7e40a40ac09ce3996e7c06873d430c40dbec21f8e4b", "fetched_at": "2026-08-28T04:10:28.607821+00:00", "kind": "readme", "missing": false, "url": "https://github.com/QwenAudio/SenseVoice"}, {"content_hash": "00799e0fcae6da1dc6bb276b1e53a07b5b1fd10d45d41f64f0bc68e66aa6d295", "fetched_at": "2026-08-29T08:23:27.860255+00:00", "kind": "homepage", "missing": false, "url": "https://huggingface.co/spaces/FunAudioLLM/SenseVoice"}], "taxonomy_version": 1}, "member_repos": {"kind": "observed", "observed_at": "2026-08-28T04:10:28.607821+00:00", "source": "github"}, "name": {"kind": "observed", "observed_at": "2026-08-28T04:10:28.607821+00:00", "source": "github"}, "platform": {"confidence": null, "enriched_at": "2026-08-29T17:23:24.024621+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "86e30a2f1672c40137c6d7e40a40ac09ce3996e7c06873d430c40dbec21f8e4b", "fetched_at": "2026-08-28T04:10:28.607821+00:00", "kind": "readme", "missing": false, "url": "https://github.com/QwenAudio/SenseVoice"}, {"content_hash": "00799e0fcae6da1dc6bb276b1e53a07b5b1fd10d45d41f64f0bc68e66aa6d295", "fetched_at": "2026-08-29T08:23:27.860255+00:00", "kind": "homepage", "missing": false, "url": "https://huggingface.co/spaces/FunAudioLLM/SenseVoice"}], "taxonomy_version": 1}, "pushed_at": {"kind": "observed", "observed_at": "2026-08-28T04:10:28.607821+00:00", "source": "github"}, "repo": {"kind": "observed", "observed_at": "2026-08-28T04:10:28.607821+00:00", "source": "github"}, "stars": {"kind": "observed", "observed_at": "2026-08-28T04:10:28.607821+00:00", "source": "github"}, "tags": {"confidence": null, "enriched_at": "2026-08-29T17:23:24.024621+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "86e30a2f1672c40137c6d7e40a40ac09ce3996e7c06873d430c40dbec21f8e4b", "fetched_at": "2026-08-28T04:10:28.607821+00:00", "kind": "readme", "missing": false, "url": "https://github.com/QwenAudio/SenseVoice"}, {"content_hash": "00799e0fcae6da1dc6bb276b1e53a07b5b1fd10d45d41f64f0bc68e66aa6d295", "fetched_at": "2026-08-29T08:23:27.860255+00:00", "kind": "homepage", "missing": false, "url": "https://huggingface.co/spaces/FunAudioLLM/SenseVoice"}], "taxonomy_version": 1}, "topics": {"kind": "observed", "observed_at": "2026-08-28T04:10:28.607821+00:00", "source": "github"}, "urls": {"kind": "observed", "observed_at": "2026-08-28T04:10:28.607821+00:00", "source": "github"}, "use_cases": {"confidence": null, "enriched_at": "2026-08-29T17:23:24.024621+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "86e30a2f1672c40137c6d7e40a40ac09ce3996e7c06873d430c40dbec21f8e4b", "fetched_at": "2026-08-28T04:10:28.607821+00:00", "kind": "readme", "missing": false, "url": "https://github.com/QwenAudio/SenseVoice"}, {"content_hash": "00799e0fcae6da1dc6bb276b1e53a07b5b1fd10d45d41f64f0bc68e66aa6d295", "fetched_at": "2026-08-29T08:23:27.860255+00:00", "kind": "homepage", "missing": false, "url": "https://huggingface.co/spaces/FunAudioLLM/SenseVoice"}], "taxonomy_version": 1}, "what_it_is": {"confidence": null, "enriched_at": "2026-08-29T17:23:24.024621+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "86e30a2f1672c40137c6d7e40a40ac09ce3996e7c06873d430c40dbec21f8e4b", "fetched_at": "2026-08-28T04:10:28.607821+00:00", "kind": "readme", "missing": false, "url": "https://github.com/QwenAudio/SenseVoice"}, {"content_hash": "00799e0fcae6da1dc6bb276b1e53a07b5b1fd10d45d41f64f0bc68e66aa6d295", "fetched_at": "2026-08-29T08:23:27.860255+00:00", "kind": "homepage", "missing": false, "url": "https://huggingface.co/spaces/FunAudioLLM/SenseVoice"}], "taxonomy_version": 1}, "when_to_avoid": {"confidence": null, "enriched_at": "2026-08-29T17:23:24.024621+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "86e30a2f1672c40137c6d7e40a40ac09ce3996e7c06873d430c40dbec21f8e4b", "fetched_at": "2026-08-28T04:10:28.607821+00:00", "kind": "readme", "missing": false, "url": "https://github.com/QwenAudio/SenseVoice"}, {"content_hash": "00799e0fcae6da1dc6bb276b1e53a07b5b1fd10d45d41f64f0bc68e66aa6d295", "fetched_at": "2026-08-29T08:23:27.860255+00:00", "kind": "homepage", "missing": false, "url": "https://huggingface.co/spaces/FunAudioLLM/SenseVoice"}], "taxonomy_version": 1}, "when_to_choose": {"confidence": null, "enriched_at": "2026-08-29T17:23:24.024621+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "86e30a2f1672c40137c6d7e40a40ac09ce3996e7c06873d430c40dbec21f8e4b", "fetched_at": "2026-08-28T04:10:28.607821+00:00", "kind": "readme", "missing": false, "url": "https://github.com/QwenAudio/SenseVoice"}, {"content_hash": "00799e0fcae6da1dc6bb276b1e53a07b5b1fd10d45d41f64f0bc68e66aa6d295", "fetched_at": "2026-08-29T08:23:27.860255+00:00", "kind": "homepage", "missing": false, "url": "https://huggingface.co/spaces/FunAudioLLM/SenseVoice"}], "taxonomy_version": 1}}, "score": {"components": {"activity": 98, "longevity": 56, "rhythm": 94}, "computed_at": "2026-09-03T02:20:16.233290+00:00", "flags": [], "formula": "round(0.45*activity + 0.35*rhythm + 0.20*longevity); archived -> min(score, 10)", "inputs": {"age_days": 791, "days_push": 16, "days_rel": 40, "gap_med": 6, "n_releases_24m": 6}, "score": 88, "version": 2}, "staleness": {"enrichment_outdated": false, "low_confidence": false, "scrape_days": 9, "stale_scrape": false}}