{
  "schema_version": "1.0.0",
  "title": "The information in diagnostic tests",
  "short_title": "Diagnostic information",
  "review_type": "mapping_review",
  "methodology_profile": "systematic-review-general-v1",
  "corresponding_owner": {
    "name": "",
    "email": "",
    "orcid": null
  },
  "contributors": [],
  "affiliations": [],
  "funding": [
    {
      "source": "National Academy of Medicine",
      "award": "2026A008797",
      "role": "The manuscript states that the funder had no role in design, acquisition, analysis, interpretation, writing, or submission."
    }
  ],
  "conflicts_of_interest": "The research team will confirm the contributor list and complete its declarations of interests before formal submission.",
  "license": "CC-BY-4.0",
  "background": "Diagnostic testing changes what is known about a target condition. Sensitivity and specificity describe test performance; likelihood ratios and post-test probabilities describe an observed result. Mutual information expresses the expected reduction in uncertainty before either result is known. Earlier studies established diagnostic applications of this measure. This study assembles a broad empirical reference scale from accessible diagnostic-review data so readers can interpret a diagnostic profile at a stated starting probability and compare its information with conventional accuracy measures. Information describes learning about a binary target. Clinical value also depends on consequences, available actions, and patient preferences.",
  "review_question": "Across reproducibly accessible diagnostic accuracy reviews, how much uncertainty about a binary target condition is a test or threshold expected to resolve at a stated starting probability, and how does this information relate to conventional test-performance measures?",
  "objectives": [
    "Build an empirical reference distribution of expected diagnostic information from source-defined tests and thresholds.",
    "Compare uncertainty resolved at 5%, 20%, and 50% starting disease probability and across a 1%-50% grid.",
    "Examine agreement and differences between diagnostic information and Youden's J using source-comparable examples.",
    "Distinguish expected information before testing from the probability update and uncertainty change after each result."
  ],
  "framework": {
    "name": "PIRD",
    "population_or_domain": "Review-defined populations and settings across binary diagnostic targets; retain source-defined population and reference standard when available.",
    "intervention_or_exposure": "A source-defined diagnostic test, rule, threshold, or analysis definition.",
    "comparators": "Conventional measures from the same sensitivity and specificity; descriptive within-review contrasts where definitions are comparable.",
    "outcomes": [
      "Mutual information in bits",
      "Percentage of starting binary uncertainty resolved",
      "Result-specific probability and uncertainty changes"
    ]
  },
  "primary_outcomes": [
    "Distribution of 100 x I(D;T)/H(D) at 5%, 20%, and 50% starting probabilities, with corresponding information in bits; each pooled estimate has equal weight."
  ],
  "secondary_outcomes": [
    "Information profiles from 1% to 50% in one-percentage-point steps.",
    "Spearman correlations with Youden's J and descriptive near-tie comparisons.",
    "CEA information crossings and Ottawa ankle rule result branches.",
    "Sensitivity to zero-cell correction and the regularized direct-binomial estimator."
  ],
  "eligibility": {
    "inclusion": [
      "Diagnostic-review data with explicit study identifiers, source links, and non-negative integer TP, FP, FN, TN counts.",
      "Groups with at least five unique study identifiers, one result per identifier, and at least three studies with both disease states represented.",
      "Aggregate sensitivity plus specificity at least one in the source-defined orientation, with one canonical representation after exact row-signature deduplication.",
      "Historical Cochrane groups linked to review-author-reported pooled findings; modern Cochrane main analyses; at most one canonical main table per OSF, Zenodo, or PMC review under the recorded priority rules."
    ],
    "exclusion": [
      "Too few studies, repeated identifiers within a group, too few studies representing both disease states, or below-chance performance in the original orientation.",
      "Counts reconstructed from rounded accuracy measures, pooled-only or consequence tables, and unsupported narrative transcription.",
      "Duplicate representations and alternatives not selected by source-specific primary designation; preserve their status in the audit record."
    ],
    "study_designs": [
      "Secondary methodological synthesis of diagnostic accuracy reviews with accessible study-level cross-classification data; original designs are inherited from each source review."
    ]
  },
  "information_sources": [
    "Fixed historical and modern Cochrane sources",
    "OSF public deposits",
    "Zenodo public deposits",
    "Figshare public deposits",
    "Dryad public deposits",
    "SRDR+ public deposits",
    "PubMed Central open-access review tables"
  ],
  "search_strategies": [
    {
      "source": "Fixed historical and modern Cochrane sources",
      "platform": "Cochrane reference dataset and published review packages",
      "query": "Retrieve CL145_open_set_20181101.zip from https://zenodo.org/records/1303259 (SHA-256 f921cf6f1a9236c86f89e7c646497639f45e8ff8fcf56008c1b27422829fec4f), and the already-held published packages for CD014911 and CD015089. Parse available study-level 2x2 tables with scripts/cochrane_dta_atlas.py at the pinned commit. This is fixed-corpus acquisition, not a new exhaustive Cochrane search.",
      "status": "executed",
      "planned_or_executed_at": null,
      "filters": [],
      "limits": "63 historical reviews and two selected modern reviews. The count and checksum identify the normalized corpus manifest. The original acquisition date needs owner confirmation.",
      "result_count": 65,
      "export_checksum": "1a17f47b78c0e2ba7c2f326111e18d2a387d04780b2a36e7b7fe54a874f51321",
      "responsible": {
        "kind": "person",
        "name": "Shuhan He"
      }
    },
    {
      "source": "OSF public deposits",
      "platform": "Frozen OSF metadata inventory and public-file inspection",
      "query": "Run each recorded query separately:\n\"diagnostic test accuracy\"\n\"diagnostic accuracy\"\n\"sensitivity and specificity\"\nsensitivity specificity meta-analysis\n\"true positive\" \"false positive\"\n\"2x2\" diagnostic",
      "status": "executed",
      "planned_or_executed_at": "2026-08-17",
      "filters": [],
      "limits": "Count and checksum identify the saved normalized discovery manifest, not a new search or included-review count. Only publicly retrievable structured diagnostic review data were eligible. OSF entries here are other researchers' evidence, not registrations of this study.",
      "result_count": 1152,
      "export_checksum": "1f46eedb7af4d6d0724d55156ba98ebb92572f3b0c4378e521400d2cd634ec44",
      "responsible": {
        "kind": "person",
        "name": "Shuhan He"
      }
    },
    {
      "source": "Zenodo public deposits",
      "platform": "Zenodo public metadata API",
      "query": "Run each recorded query separately:\n\"diagnostic test accuracy\"\n\"diagnostic accuracy\" AND \"systematic review\"\n\"diagnostic accuracy\" AND \"meta-analysis\"\n\"test accuracy\" AND \"meta-analysis\"\n\"sensitivity and specificity\" AND \"meta-analysis\"\n\"true positive\" AND \"false positive\" AND \"meta-analysis\"\n\"2x2\" AND diagnostic AND \"meta-analysis\"",
      "status": "executed",
      "planned_or_executed_at": "2026-08-17",
      "filters": [],
      "limits": "Count and checksum identify the saved normalized discovery manifest, not a new search or included-review count. Only publicly retrievable structured diagnostic review data were eligible. OSF entries here are other researchers' evidence, not registrations of this study.",
      "result_count": 799,
      "export_checksum": "0fa8e301ef599db6657089e96ef341e816f790e444a2c5fe28da202d4e61a44c",
      "responsible": {
        "kind": "person",
        "name": "Shuhan He"
      }
    },
    {
      "source": "Figshare public deposits",
      "platform": "Figshare public metadata API",
      "query": "Run each recorded query separately:\n\"diagnostic test accuracy\"\n\"diagnostic accuracy\" AND \"systematic review\"\n\"diagnostic accuracy\" AND \"meta-analysis\"\n\"test accuracy\" AND \"meta-analysis\"\n\"sensitivity and specificity\" AND \"meta-analysis\"\n\"true positive\" AND \"false positive\" AND \"meta-analysis\"\n\"2x2\" AND diagnostic AND \"meta-analysis\"",
      "status": "executed",
      "planned_or_executed_at": "2026-08-17",
      "filters": [],
      "limits": "Count and checksum identify the saved normalized discovery manifest, not a new search or included-review count. Only publicly retrievable structured diagnostic review data were eligible. OSF entries here are other researchers' evidence, not registrations of this study.",
      "result_count": 262,
      "export_checksum": "b38dfd72097e51a97cc52f0d679097e3897e29249ad7a7cf11c50a14bb0efb4f",
      "responsible": {
        "kind": "person",
        "name": "Shuhan He"
      }
    },
    {
      "source": "Dryad public deposits",
      "platform": "Dryad public metadata API",
      "query": "Run each recorded query separately:\n\"diagnostic test accuracy\"\n\"diagnostic accuracy\" AND \"systematic review\"\n\"diagnostic accuracy\" AND \"meta-analysis\"\n\"test accuracy\" AND \"meta-analysis\"\n\"sensitivity and specificity\" AND \"meta-analysis\"\n\"true positive\" AND \"false positive\" AND \"meta-analysis\"\n\"2x2\" AND diagnostic AND \"meta-analysis\"",
      "status": "executed",
      "planned_or_executed_at": "2026-08-17",
      "filters": [],
      "limits": "Count and checksum identify the saved normalized discovery manifest, not a new search or included-review count. Only publicly retrievable structured diagnostic review data were eligible. OSF entries here are other researchers' evidence, not registrations of this study.",
      "result_count": 11,
      "export_checksum": "e2e5e7913d05a578b36fbaaca5f107bf5d2bb1d0b8905ad2953c43e4e1a5d52f",
      "responsible": {
        "kind": "person",
        "name": "Shuhan He"
      }
    },
    {
      "source": "SRDR+ public deposits",
      "platform": "DataCite SRDR metadata and fixed public SRDR inventory",
      "query": "publisher:\"Systematic Review Data Repository\"; retain the fixed inventory of 172 projects and apply the source-specific discovery and diagnostic-candidate rules in scripts/open_repository_dta.py.",
      "status": "executed",
      "planned_or_executed_at": "2026-08-17",
      "filters": [],
      "limits": "Count and checksum identify the saved normalized discovery manifest, not a new search or included-review count. Only publicly retrievable structured diagnostic review data were eligible. OSF entries here are other researchers' evidence, not registrations of this study.",
      "result_count": 172,
      "export_checksum": "0576f936db6c27b7cde5147f346289b08c399d7ef90e5f96902a066c669ff88a",
      "responsible": {
        "kind": "person",
        "name": "Shuhan He"
      }
    },
    {
      "source": "PubMed Central open-access review tables",
      "platform": "NCBI E-Utilities and PMC JATS XML",
      "query": "((\"diagnostic test accuracy\"[Title/Abstract]) OR (\"diagnostic accuracy\"[Title/Abstract])) AND (\"2x2\"[All Fields] OR \"2 x 2\"[All Fields] OR (\"true positive\"[All Fields] AND \"false positive\"[All Fields] AND \"false negative\"[All Fields] AND \"true negative\"[All Fields])) AND (systematic review[Title/Abstract] OR meta-analysis[Title/Abstract]) AND open access[filter] AND 2020:2026[Publication Date] NOT protocol[Title]",
      "status": "executed",
      "planned_or_executed_at": "2026-08-17",
      "filters": [
        "open access",
        "2020:2026 publication dates",
        "exclude protocol in title"
      ],
      "limits": "The fixed query returned 1,518 articles on 2026-08-17; all were retrieved. Only explicit identifiers and integer TP, FP, FN, TN columns were normalized. Count and checksum identify articles.csv.",
      "result_count": 1518,
      "export_checksum": "c0dbf4a9a700ef837e76b1f8551a1cb09c237d93a1ce1c8b9465ff731b81f8f3",
      "responsible": {
        "kind": "person",
        "name": "Shuhan He"
      }
    }
  ],
  "date_limits": {
    "from": null,
    "to": null,
    "justification": "No common date window: historical Cochrane is a fixed 2018 release; PMC uses 2020-2026 publication years; other sources use recorded inventories or queries. Source-specific limits govern inclusion."
  },
  "language_limits": {
    "languages": [],
    "justification": "No separate language restriction is documented in the inspected methods. Query terms are English; this empty list does not establish language-complete coverage."
  },
  "grey_literature_plan": "Public OSF, Zenodo, Figshare, Dryad, and SRDR deposits were examined for reusable diagnostic-review data. Dryad and SRDR contributed no primary pooled estimates. This accessible evidence sample is not an exhaustive search of unpublished studies or all diagnostic medicine.",
  "deduplication_method": "Retain source identifiers and compare exact normalized count-row signatures. Keep one result per identifier within a group; repeated identifiers make that group ineligible. The same study may contribute to different diagnostic profiles, so study-result appearances are not unique studies.",
  "screening_method": {
    "independent_screeners": 1,
    "conflict_resolution": "One author-directed computational selection workflow and a recorded author audit are documented. Two independent human screeners are not established. Fixed source priorities resolved representations; ambiguous groups were excluded or retained outside the primary sample. The owner must confirm this retrospective description.",
    "ai_role": "Codex assisted retrieval, parsing and code drafting. Deterministic rules selected profiles; agent outputs do not constitute independent human screening."
  },
  "full_text_method": "Inspect structured review deposits and machine-readable PMC tables for explicit study-level counts. Preserve article, table, file and review identifiers and exclusion reasons. No complete two-reviewer full-text screening process is claimed.",
  "data_extraction_method": "Parse explicit study identifiers and integer TP, FP, FN, TN cells while retaining source location, orientation and checksums. Do not reconstruct counts from rounded accuracy metrics. Apply canonical workbook/table rules. Compare fitted estimates with usable source summaries and retain the recorded author audit; that audit is not independent adjudication.",
  "risk_of_bias_method": "No new uniform study-level QUADAS-2 assessment across the corpus is documented. Incomplete source risk-of-bias and clinical-context reporting limit interpretation. This record does not invent appraisal judgments or equate numerical model checks with risk-of-bias assessment.",
  "synthesis_plan": "Pool each eligible source-defined test or threshold separately with bivariate random-effects REML on sensitivity/specificity logits, adding 0.5 to all four cells in every study. Back-transform the fitted mean logits. Compute entropy, mutual information and percentage uncertainty resolved at the specified starting probabilities. Give each pooled estimate equal reference-distribution weight. Report pooled-mean intervals and plug-in predictive ranges conditional on fitted heterogeneity. Bootstrap reference medians over pooled estimates with 2,000 percentile replicates and a fixed seed; this does not propagate within-estimate uncertainty. Repeat analyses with correction only in zero-containing studies and with the regularized direct-binomial estimator. Treat within-review contrasts and CEA crossings as descriptive; no paired superiority inference or clinical utility ranking is supported. Distinguish signed entropy change, posterior-to-prior KL divergence, log2 likelihood ratio, and expected mutual information.",
  "certainty_assessment": "No new GRADE ratings are claimed for this methodological reference scale. Report source fidelity, numerical diagnostics, estimator sensitivity and sample limitations separately from certainty about clinical effects or recommendations.",
  "reporting_guidelines": [
    "PRISMA-P",
    "The BMJ Research Methods and Reporting guidance"
  ],
  "anticipated_start_date": "",
  "anticipated_completion_date": "",
  "milestones": {
    "pilot_search_started_at": null,
    "formal_search_started_at": null,
    "screening_started_at": null,
    "extraction_started_at": null,
    "synthesis_started_at": "2026-08-16"
  },
  "timing_explanation": "Retrospective draft prepared on 2026-09-08 after analysis and manuscript preparation. Commit 857cd0f on 2026-08-16 documents an executed Cochrane atlas synthesis. The synthesis field records this evidenced run as a provisional timing anchor, not a confirmed first-ever analysis date. Formal search, screening, extraction and original planned dates remain unconfirmed and unfilled. Repository/PMC expansion runs are dated 2026-08-17; manuscript 1.10.1 was committed on 2026-08-30. This record is not prospective. Primary-designation rules developed during staged audits were fixed before the final reference analysis; this does not establish that they preceded all data inspection. The owner must confirm actual milestones before submission. PRISMA-P organizes this retrospective record and does not imply prior adherence or registration.",
  "automation": {
    "agent_name": "OpenAI Codex",
    "agent_version": null,
    "provider": "OpenAI",
    "roles": [
      "Retrospective protocol drafting from versioned sources",
      "Research-use code drafting, retrieval, deterministic checks and language editing as disclosed in the manuscript"
    ],
    "human_reviewed_fields": [],
    "source_artifacts": [
      "https://github.com/ShuhanCS/diagnostic-information/blob/0c8cff2d06233e0fb7dad4e3067e69c15ca6e5c7/docs/manuscript/releases/v1.10.1/sources/draft-v1.10.1.md",
      "https://github.com/ShuhanCS/diagnostic-information/blob/0c8cff2d06233e0fb7dad4e3067e69c15ca6e5c7/docs/manuscript/releases/v1.10.1/sources/supplement-v1.10.1.md",
      "https://github.com/ShuhanCS/diagnostic-information/blob/0c8cff2d06233e0fb7dad4e3067e69c15ca6e5c7/docs/manuscript/AUTHORSHIP.json",
      "https://github.com/ShuhanCS/diagnostic-information/blob/0c8cff2d06233e0fb7dad4e3067e69c15ca6e5c7/docs/specs/2026-08-16-full-cochrane-dta-atlas-spec.md",
      "https://github.com/ShuhanCS/diagnostic-information/blob/0c8cff2d06233e0fb7dad4e3067e69c15ca6e5c7/docs/manuscript/releases/v1.10.1/MANIFEST.json",
      "https://github.com/ShuhanCS/diagnostic-information/blob/0c8cff2d06233e0fb7dad4e3067e69c15ca6e5c7/scripts/open_repository_dta.py",
      "https://github.com/ShuhanCS/diagnostic-information/blob/0c8cff2d06233e0fb7dad4e3067e69c15ca6e5c7/scripts/pmc_pilot.py",
      "https://github.com/ShuhanCS/diagnostic-information/blob/0c8cff2d06233e0fb7dad4e3067e69c15ca6e5c7/analysis/cochrane_atlas/corpus_manifest.csv",
      "https://github.com/ShuhanCS/diagnostic-information/blob/0c8cff2d06233e0fb7dad4e3067e69c15ca6e5c7/analysis/open_repository_atlas/osf/source_manifest.csv",
      "https://github.com/ShuhanCS/diagnostic-information/blob/0c8cff2d06233e0fb7dad4e3067e69c15ca6e5c7/analysis/open_repository_atlas/zenodo/source_manifest.csv",
      "https://github.com/ShuhanCS/diagnostic-information/blob/0c8cff2d06233e0fb7dad4e3067e69c15ca6e5c7/analysis/open_repository_atlas/figshare/source_manifest.csv",
      "https://github.com/ShuhanCS/diagnostic-information/blob/0c8cff2d06233e0fb7dad4e3067e69c15ca6e5c7/analysis/open_repository_atlas/dryad/source_manifest.csv",
      "https://github.com/ShuhanCS/diagnostic-information/blob/0c8cff2d06233e0fb7dad4e3067e69c15ca6e5c7/analysis/open_repository_atlas/srdr/source_manifest.csv",
      "https://github.com/ShuhanCS/diagnostic-information/blob/0c8cff2d06233e0fb7dad4e3067e69c15ca6e5c7/analysis/pmc_open_atlas/articles.csv"
    ],
    "limitations": "This new record has not received owner or coauthor approval. The manuscript reports SH review of earlier AI-assisted work; that does not approve this protocol. Source repositories are private unless access is granted. The proposed CC-BY-4.0 license applies only to protocol text after owner acceptance, not third-party data. Cochrane data retain their original restrictions. Unfilled dates require confirmation. No registration receipt, DOI, independent peer review, or prospective registration is claimed. Draft record 1.0.2 leaves the author list in preparation following the owner's correction. Contact details are supplied privately from the signed-in Publishing account. Methods and unconfirmed dates remain as prepared."
  }
}
