{
  "@context": "https://schema.org",
  "@type": "Claim",
  "@id": "https://wulfkaal.github.io/claims/7261018-002",
  "identifier": "kaal:claim:7261018-002",
  "text": "What they do not measure, because their designs contain no mechanism by which an agent’s payoff depends on the verified quality of its work, is: whether the agents’ reports about their work track the work itself; whether agents evaluate one another independently or herd on the visible consensus; whether confident answers are calibrated answers; whether agents contribute to collective evaluation or free-ride on it.",
  "author": {
    "@type": "Person",
    "name": "Wulf A. Kaal",
    "identifier": "https://orcid.org/0009-0008-7840-1847"
  },
  "citation": "Wulf A. Kaal, Empirical Evaluation of the Agentic Reputation Substrate: Deliberation, the Composition of Error, and the Registered Measurement of Agency Costs in a Controlled Multi-Model Cohort (2026). SSRN: https://ssrn.com/abstract=7261018",
  "isBasedOn": {
    "@type": "ScholarlyArticle",
    "name": "Empirical Evaluation of the Agentic Reputation Substrate: Deliberation, the Composition of Error, and the Registered Measurement of Agency Costs in a Controlled Multi-Model Cohort",
    "datePublished": "2026",
    "url": "https://ssrn.com/abstract=7261018",
    "sha256": "1d6cbe544bd0055133f7cf8ff308be4fde955867bd8dc764992b8d516af15fa8",
    "encoding": "https://raw.githubusercontent.com/wulfkaal/Academic-Papers/main/papers/pdf/Kaal%20-%202026%20-%20Empirical%20Evaluation%20of%20the%20Agentic%20Reputation%20Substrate%20Deliberation%2C%20the%20Composition%20of%20Error%2C%20and%20the%20Registered%20Measurement%20of%20Agency%20Costs%20in%20a%20Controlled%20Multi-Model%20Cohort.pdf"
  },
  "keywords": [
    "ai-and-agents",
    "economics",
    "risk-and-incentives"
  ],
  "about": [],
  "abstract": "What they do not measure, because their designs contain no mechanism by which an agent’s payoff depends on the verified quality of its work, is: whether the agents’ reports about their work track the work itself; whether agents evaluate one another independently or herd on the visible consensus; whether confident answers are calibrated answers; whether agents contribute to collective evaluation or free-ride on it.",
  "additionalProperty": [
    {
      "@type": "PropertyValue",
      "name": "claim_type",
      "value": "condition"
    },
    {
      "@type": "PropertyValue",
      "name": "confidence",
      "value": "argued"
    },
    {
      "@type": "PropertyValue",
      "name": "scope_conditions",
      "value": [
        "benchmark designs in which verified work quality does not affect agent payoff"
      ]
    },
    {
      "@type": "PropertyValue",
      "name": "is_failure_mode",
      "value": false
    },
    {
      "@type": "PropertyValue",
      "name": "edges",
      "value": []
    }
  ],
  "isPartOf": {
    "@id": "https://wulfkaal.github.io/claims/index.json"
  },
  "dateModified": "2026-08-10",
  "version": "1.0",
  "sha256": "04b84984116c1f7eb22a722ec119a0e581302e95e9171868dfe8ae518941f6b7",
  "canonicalForm": "https://wulfkaal.github.io/claims/7261018-002.md"
}
