{
 "@context": "https://schema.org",
 "@type": "Claim",
 "@id": "https://wulfkaal.github.io/positions/2026-08-08-332",
 "identifier": "kaal:position:2026-08-08-332",
 "additionalType": "https://wulfkaal.github.io/positions/schema.json#AffirmedPositionClaim",
 "name": "Pooled Confidence Does Not Establish Cross Setting Prediction",
 "text": "IntHout, Ioannidis, Rovers, and Goeman qualify the inference from a pooled confidence interval. Under between-study heterogeneity, a pooled 95 percent confidence interval may remain wholly on one side of zero while a prediction interval reaches both sides. Their review found that 72.4 percent of statistically significant Cochrane meta-analyses with observed heterogeneity had prediction intervals that included the null. This result does not dispute the registered calculation or the negative direction in each campaign unit. It shows that pooled significance and directional consistency do not alone establish how the effect will generalize to another exchangeable unit. That stronger claim requires a heterogeneity-aware prediction analysis. The paper does not test Kaal's cohort, hypotheses, estimator, or campaign units.",
 "author": {
  "@type": "Person",
  "name": "Wulf A. Kaal",
  "identifier": "https://orcid.org/0009-0008-7840-1847"
 },
 "datePublished": "2026-08-08",
 "dateModified": "2026-08-08",
 "creativeWorkStatus": "Affirmed",
 "responseType": "qualification",
 "keywords": [
  "research-methods",
  "consensus-and-security",
  "ai-and-agents",
  "scholarly-growth-coverage",
  "scholarly-literature",
  "statistical-inference",
  "confidence-intervals",
  "meta-analysis",
  "heterogeneity",
  "prediction-intervals",
  "evidence-provenance"
 ],
 "scope_conditions": [
  "The response is limited to the exact BMJ Open passages and the one mapped Kaal claim.",
  "External evidence level: complete public BMJ Open article in Europe PMC JATS XML with a concordant Crossref record and exact proposition-bearing passages.",
  "Mapping review tier: independent substantive scholarly-growth qualification.",
  "The paper studies clinical meta-analyses and does not evaluate Kaal's multi-model cohort, registered hypotheses, estimator, or campaign units.",
  "The source does not recalculate Kaal's pooled intervals or estimate between-campaign heterogeneity.",
  "Every observed campaign contrast being negative is stronger descriptive evidence than a pooled confidence interval alone, but it is not a prediction interval for another exchangeable campaign unit.",
  "The qualification addresses generalization only. It does not dispute the registered calculation or its reported directional consistency.",
  "Crossref Spanish and Semantic Scholar search returned HTTP 429. Several inherited identity refreshes were rate-limited or unresolved. No blocked response was promoted."
 ],
 "currentDebate": {
  "name": "Plea for routinely presenting prediction intervals in meta-analysis",
  "url": "https://doi.org/10.1136/bmjopen-2015-010247"
 },
 "extends": {
  "identifier": "kaal:claim:7261018-036",
  "url": "https://wulfkaal.github.io/claims/7261018-036",
  "citation": "Wulf A. Kaal, Empirical Evaluation of the Agentic Reputation Substrate: Deliberation, the Composition of Error, and the Registered Measurement of Agency Costs in a Controlled Multi-Model Cohort (2026). SSRN: https://ssrn.com/abstract=7261018",
  "paper": "Wulf A. Kaal, Empirical Evaluation of the Agentic Reputation Substrate: Deliberation, the Composition of Error, and the Registered Measurement of Agency Costs in a Controlled Multi-Model Cohort",
  "authors": [
   "Wulf A. Kaal"
  ],
  "year": "2026",
  "ssrn": "https://ssrn.com/abstract=7261018",
  "source_pdf_sha256": "1d6cbe544bd0055133f7cf8ff308be4fde955867bd8dc764992b8d516af15fa8"
 },
 "isBasedOn": [
  {
   "@id": "https://wulfkaal.github.io/claims/7261018-036"
  },
  {
   "@type": "CreativeWork",
   "name": "Plea for routinely presenting prediction intervals in meta-analysis",
   "url": "https://doi.org/10.1136/bmjopen-2015-010247"
  }
 ],
 "batch_id": "kaal-review:2026-08-12:scholarly-growth-7261018-036-reviewed-v1",
 "review_provenance": "https://wulfkaal.github.io/positions/by-claim/7261018-036.html",
 "publicationStatus": "public",
 "recordTypeNote": "Dated commentary position extending a scholarly corpus claim. Not a verbatim claim extracted from the paper.",
 "isPartOf": {
  "@id": "https://wulfkaal.github.io/positions/index.json"
 },
 "version": "1.0",
 "canonical_url": "https://wulfkaal.github.io/positions/2026-08-08-332",
 "canonicalForm": "https://wulfkaal.github.io/positions/2026-08-08-332.md",
 "candidateId": "kaal:response-candidate:2026-08-12:scholarly-growth-7261018-036-inthout-et-al-01",
 "evidenceLevel": "complete public BMJ Open article in Europe PMC JATS XML with a concordant Crossref record and exact proposition-bearing passages",
 "reviewTier": "independent substantive scholarly-growth qualification",
 "mappingConfidence": 0.98,
 "mappingAmbiguous": false,
 "mappingMethod": "independent substantive scholarly-growth one-to-one review",
 "mappingWhyRelevant": "The paper directly distinguishes inference about a pooled average confidence interval from prediction across heterogeneous independent settings.",
 "sourceProvenance": {
  "source": "BMJ Open article with complete Europe PMC full text and concordant Crossref identity",
  "sourceRecordId": "doi:10.1136/bmjopen-2015-010247",
  "canonicalUrl": "https://doi.org/10.1136/bmjopen-2015-010247",
  "landingPageUrl": "https://bmjopen.bmj.com/content/6/7/e010247",
  "fullTextUrl": "https://pmc.ncbi.nlm.nih.gov/articles/PMC4947751/",
  "retrievedAt": "2026-08-13T04:39:26.818Z",
  "crossrefRecordSha256": "b89bc6015e248663e2306a2f489d628fc76f95509110ca409fb8e9e488115985",
  "officialLandingPageStatus": 403,
  "publicArticlePageSha256": "599bb39663bb29d7f85a2f80de166019d4cd14bacea97eb5c75a674abfd89c67",
  "fullTextSha256": "11664ef42f87609445b8fbad866095deb33c7e950a2d99498cc61b8a7572d07e",
  "textExtraction": {
   "tool": "JATS XML tag normalization",
   "quality": "complete readable public scholarly text with exact methods and results passages"
  },
  "sourceProposition": "IntHout and coauthors show that a pooled 95 percent confidence interval can remain wholly on one side of the null while a prediction interval spans both sides when between-study heterogeneity is present. A prediction interval addresses the range of true effects expected across similar settings, not only the average pooled effect.",
  "sourcePropositionSha256": "fefd69d0aa11f8a0598cf37ecc5a3cb0429759b66f04f97d8b975e486ff6f290",
  "sourceEvidenceSetSha256": "df7f6cdd505d73f7b0055d8c4956e4d3a14d63810c7b57e3546c1edd0fa7439c",
  "sourceEvidencePassages": [
   {
    "text": "A 95% prediction interval estimates where the true effects are to be expected for 95% of similar (exchangeable) studies that might be conducted in the future.",
    "locator": {
     "source": "BMJ Open 6:e010247 (2016)",
     "section": "Methods: Prediction intervals",
     "fullTextFormat": "Europe PMC JATS XML",
     "fullTextSha256": "11664ef42f87609445b8fbad866095deb33c7e950a2d99498cc61b8a7572d07e"
    },
    "sha256": "b009918a71efd519dc945fcdbd86887a4a7ddb7f6ef8953f56c3372cdbdebbad"
   },
   {
    "text": "However, in case of heterogeneity, a prediction interval covers a wider range than a CI. Consequently, in case of a statistically significant effect (where all values of the 95% CI are on the same side of the null), the corresponding 95% prediction interval may indicate that values are possible on both sides of the null.",
    "locator": {
     "source": "BMJ Open 6:e010247 (2016)",
     "section": "Methods: Prediction intervals",
     "fullTextFormat": "Europe PMC JATS XML",
     "fullTextSha256": "11664ef42f87609445b8fbad866095deb33c7e950a2d99498cc61b8a7572d07e"
    },
    "sha256": "a410a7793a6aa952a27a897662aac33daebdbae89e832cedd5f5155be3b17e6a"
   },
   {
    "text": "Consequently, almost three-quarter (347, 72.4%) had a prediction interval that contained the null effect.",
    "locator": {
     "source": "BMJ Open 6:e010247 (2016)",
     "section": "Results: Prediction intervals containing the null effect",
     "fullTextFormat": "Europe PMC JATS XML",
     "fullTextSha256": "11664ef42f87609445b8fbad866095deb33c7e950a2d99498cc61b8a7572d07e"
    },
    "sha256": "a3ce35829a4d9cfe29343fe71a581d7aa508efe01ef6823e0c73f2ef84dda578"
   }
  ],
  "workId": "work:doi:10.1136/bmjopen-2015-010247",
  "workAuthors": [
   "Joanna IntHout",
   "John P A Ioannidis",
   "Maroeska M Rovers",
   "Jelle J Goeman"
  ],
  "workPublishedAt": "2016-07-12",
  "identityKeys": [
   "doi:10.1136/bmjopen-2015-010247",
   "pmcid:PMC4947751",
   "crossref:b89bc6015e248663e2306a2f489d628fc76f95509110ca409fb8e9e488115985",
   "proposition:fefd69d0aa11f8a0598cf37ecc5a3cb0429759b66f04f97d8b975e486ff6f290"
  ],
  "claimMappings": [
   {
    "claimId": "kaal:claim:7261018-036",
    "claimUrl": "https://wulfkaal.github.io/claims/7261018-036",
    "rank": 1,
    "confidence": 0.98,
    "method": "independent substantive scholarly-growth one-to-one review",
    "whyRelevant": "The paper directly distinguishes inference about a pooled average confidence interval from prediction across heterogeneous independent settings.",
    "ambiguous": false
   }
  ],
  "substantiveReview": {
   "reviewedAt": "2026-08-13T04:43:09.015Z",
   "sourceIdentityVerified": true,
   "authorIndependenceVerified": true,
   "canonicalPublicStatusVerified": true,
   "retractionOrSupersessionFound": false,
   "propositionFidelityVerified": true,
   "mechanismCorrespondence": "pooled confidence interval inference compared with prediction across heterogeneous independent settings",
   "compatibleScope": "external methodological qualification limited to generalization beyond the pooled average and observed campaign units",
   "responseWordingDefensible": true,
   "oneToOneExtendsMapping": true,
   "sameSourceDistinctionVerified": true,
   "exactSupportingQuotesVerified": true,
   "limitations": [
    "The paper studies clinical meta-analyses and does not evaluate Kaal's multi-model cohort, registered hypotheses, estimator, or campaign units.",
    "The source does not recalculate Kaal's pooled intervals or estimate between-campaign heterogeneity.",
    "Every observed campaign contrast being negative is stronger descriptive evidence than a pooled confidence interval alone, but it is not a prediction interval for another exchangeable campaign unit.",
    "The qualification addresses generalization only. It does not dispute the registered calculation or its reported directional consistency.",
    "Crossref Spanish and Semantic Scholar search returned HTTP 429. Several inherited identity refreshes were rate-limited or unresolved. No blocked response was promoted."
   ]
  },
  "liveVerification": {
   "checkedAt": "2026-08-13T04:44:28.631Z",
   "crossrefHttpStatus": 200,
   "crossrefResponseSha256": "b89bc6015e248663e2306a2f489d628fc76f95509110ca409fb8e9e488115985",
   "fullTextHttpStatus": 200,
   "fullTextSha256": "11664ef42f87609445b8fbad866095deb33c7e950a2d99498cc61b8a7572d07e",
   "frozenFullTextSha256": "11664ef42f87609445b8fbad866095deb33c7e950a2d99498cc61b8a7572d07e"
  }
 },
 "userAffirmation": "Automatically authorized under standing authority receipt kaal-standing-publication-authorization:2026-08-01:hourly-reviewed-batches, SHA-256 e2126054b58bb4e88db65c334ef4fc8ae78dcaadc63543c408133c4eefceb9b1, limited to this substantively reviewed scholarly-growth qualification and the protected scholarly claims matching the current bridge checkpoint under the current owner instruction.",
 "sha256": "27e0dd506806c5d2f46cdba6d0336126ea8fb7c6557f63d116188511a2cfd66f"
}
