{"adoption": {"forks": 312, "observed_at": "2026-08-28T04:03:24.744949+00:00", "stars": 1057}, "canonical_url": "https://ross.abutalabs.com/products/adam", "card": {"archived": false, "artifact_type": "library", "description": "ADAM is a genomics analysis platform with specialized file formats built using Apache Avro, Apache Spark, and Apache Parquet. Apache 2 licensed.", "domain": ["bioinformatics", "big-data"], "enriched": true, "function": ["etl", "data-science", "cli", "serialization"], "health_score": 64, "homepage": null, "language": "Scala", "license": "Apache-2.0", "license_family": "permissive", "maturity": "active", "member_repos": ["bigdatagenomics/adam"], "name": "bigdatagenomics/adam", "platform": ["jvm", "python", "go", "rust", "cli", "cloud"], "pushed_at": "2026-03-17T20:10:53+00:00", "repo": "bigdatagenomics/adam", "stars": 1057, "tags": ["genomics", "apache-spark", "parquet", "avro", "variant-calling", "sequencing", "scala", "distributed-computing", "data-engineering"], "topics": ["spark", "big-data", "bioinformatics", "genomics", "parquet", "avro", "scala", "java", "python", "r"], "urls": [], "use_cases": ["run DNA sequencing pipelines on a Spark cluster", "convert BAM and VCF files to Parquet for faster analytics", "parallelize variant calling across thousands of cores", "process RNA-seq read counts at scale", "query genomic data with SQL from Spark", "replace script-glued sequencing workflows with an in-memory pipeline"], "what_it_is": "ADAM is a genomics analysis platform built on Apache Spark that provides schemas and APIs for processing genomic data like reads, variants, and features. It works as both a Scala/Java/Python/R library and a CLI tool, using Avro and Parquet-based formats with support for legacy formats like BAM, VCF, and BED.", "when_to_avoid": ["you only need small-scale, single-sample analysis where standard tools like GATK or samtools suffice", "you cannot operate a Spark cluster or JVM environment", "you need a turnkey GUI-based genomics platform"], "when_to_choose": ["you need to scale genomics processing beyond a single machine", "you want interoperable columnar storage for genomic data", "your team already uses Spark and wants genomic APIs in Scala, Java, Python, R, or SQL"]}, "data_as_of": "2026-08-30T08:39:29.467469+00:00", "members": [{"path": "/products/adam", "repo": "bigdatagenomics/adam", "role": "main", "score": 55}], "provenance": {"archived": {"kind": "observed", "observed_at": "2026-08-28T04:03:24.744949+00:00", "source": "github"}, "artifact_type": {"confidence": null, "enriched_at": "2026-08-30T06:57:56.077460+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "6d11f59cb7d5f20472537ba3e1fa883b868c57b8848885fdb6abe7383669ecbf", "fetched_at": "2026-08-28T04:03:24.744949+00:00", "kind": "readme", "missing": false, "url": "https://github.com/bigdatagenomics/adam"}], "taxonomy_version": 1}, "description": {"kind": "observed", "observed_at": "2026-08-28T04:03:24.744949+00:00", "source": "github"}, "domain": {"confidence": null, "enriched_at": "2026-08-30T06:57:56.077460+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "6d11f59cb7d5f20472537ba3e1fa883b868c57b8848885fdb6abe7383669ecbf", "fetched_at": "2026-08-28T04:03:24.744949+00:00", "kind": "readme", "missing": false, "url": "https://github.com/bigdatagenomics/adam"}], "taxonomy_version": 1}, "enriched": {"inputs": [], "kind": "computed", "method": "enrichment_status"}, "function": {"confidence": null, "enriched_at": "2026-08-30T06:57:56.077460+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "6d11f59cb7d5f20472537ba3e1fa883b868c57b8848885fdb6abe7383669ecbf", "fetched_at": "2026-08-28T04:03:24.744949+00:00", "kind": "readme", "missing": false, "url": "https://github.com/bigdatagenomics/adam"}], "taxonomy_version": 1}, "health_score": {"inputs": ["days_since_push", "days_since_release", "archived"], "kind": "computed", "method": "health_v1"}, "homepage": {"kind": "observed", "observed_at": "2026-08-28T04:03:24.744949+00:00", "source": "github"}, "language": {"kind": "observed", "observed_at": "2026-08-28T04:03:24.744949+00:00", "source": "github"}, "license": {"kind": "observed", "observed_at": "2026-08-28T04:03:24.744949+00:00", "source": "github"}, "license_family": {"inputs": ["license"], "kind": "computed", "method": "license_family"}, "maturity": {"confidence": null, "enriched_at": "2026-08-30T06:57:56.077460+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "6d11f59cb7d5f20472537ba3e1fa883b868c57b8848885fdb6abe7383669ecbf", "fetched_at": "2026-08-28T04:03:24.744949+00:00", "kind": "readme", "missing": false, "url": "https://github.com/bigdatagenomics/adam"}], "taxonomy_version": 1}, "member_repos": {"kind": "observed", "observed_at": "2026-08-28T04:03:24.744949+00:00", "source": "github"}, "name": {"kind": "observed", "observed_at": "2026-08-28T04:03:24.744949+00:00", "source": "github"}, "platform": {"confidence": null, "enriched_at": "2026-08-30T06:57:56.077460+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "6d11f59cb7d5f20472537ba3e1fa883b868c57b8848885fdb6abe7383669ecbf", "fetched_at": "2026-08-28T04:03:24.744949+00:00", "kind": "readme", "missing": false, "url": "https://github.com/bigdatagenomics/adam"}], "taxonomy_version": 1}, "pushed_at": {"kind": "observed", "observed_at": "2026-08-28T04:03:24.744949+00:00", "source": "github"}, "repo": {"kind": "observed", "observed_at": "2026-08-28T04:03:24.744949+00:00", "source": "github"}, "stars": {"kind": "observed", "observed_at": "2026-08-28T04:03:24.744949+00:00", "source": "github"}, "tags": {"confidence": null, "enriched_at": "2026-08-30T06:57:56.077460+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "6d11f59cb7d5f20472537ba3e1fa883b868c57b8848885fdb6abe7383669ecbf", "fetched_at": "2026-08-28T04:03:24.744949+00:00", "kind": "readme", "missing": false, "url": "https://github.com/bigdatagenomics/adam"}], "taxonomy_version": 1}, "topics": {"kind": "observed", "observed_at": "2026-08-28T04:03:24.744949+00:00", "source": "github"}, "urls": {"kind": "observed", "observed_at": "2026-08-28T04:03:24.744949+00:00", "source": "github"}, "use_cases": {"confidence": null, "enriched_at": "2026-08-30T06:57:56.077460+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "6d11f59cb7d5f20472537ba3e1fa883b868c57b8848885fdb6abe7383669ecbf", "fetched_at": "2026-08-28T04:03:24.744949+00:00", "kind": "readme", "missing": false, "url": "https://github.com/bigdatagenomics/adam"}], "taxonomy_version": 1}, "what_it_is": {"confidence": null, "enriched_at": "2026-08-30T06:57:56.077460+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "6d11f59cb7d5f20472537ba3e1fa883b868c57b8848885fdb6abe7383669ecbf", "fetched_at": "2026-08-28T04:03:24.744949+00:00", "kind": "readme", "missing": false, "url": "https://github.com/bigdatagenomics/adam"}], "taxonomy_version": 1}, "when_to_avoid": {"confidence": null, "enriched_at": "2026-08-30T06:57:56.077460+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "6d11f59cb7d5f20472537ba3e1fa883b868c57b8848885fdb6abe7383669ecbf", "fetched_at": "2026-08-28T04:03:24.744949+00:00", "kind": "readme", "missing": false, "url": "https://github.com/bigdatagenomics/adam"}], "taxonomy_version": 1}, "when_to_choose": {"confidence": null, "enriched_at": "2026-08-30T06:57:56.077460+00:00", "kind": "inferred", "prompt_version": 1, "sources": [{"content_hash": "6d11f59cb7d5f20472537ba3e1fa883b868c57b8848885fdb6abe7383669ecbf", "fetched_at": "2026-08-28T04:03:24.744949+00:00", "kind": "readme", "missing": false, "url": "https://github.com/bigdatagenomics/adam"}], "taxonomy_version": 1}}, "score": {"components": {"activity": 72, "longevity": 100, "rhythm": 8}, "computed_at": "2026-09-03T02:20:16.233290+00:00", "flags": [], "formula": "round(0.45*activity + 0.35*rhythm + 0.20*longevity); archived -> min(score, 10)", "inputs": {"age_days": 4670, "days_push": 169, "days_rel": null, "gap_med": null, "n_releases_24m": 0}, "score": 55, "version": 2}, "staleness": {"enrichment_outdated": false, "low_confidence": false, "scrape_days": 9, "stale_scrape": false}}