{
 "@context": "https://schema.org",
 "@type": "Claim",
 "@id": "https://wulfkaal.github.io/claims/4855607-014",
 "identifier": "kaal:claim:4855607-014",
 "text": "Deep reinforcement learning demands large amounts of training data, which suggests its algorithms differ fundamentally from human learning, and learning without supervision becomes particularly hard when rewards are sparse, as they typically are in sequence generation tasks.",
 "author": {
  "@type": "Person",
  "name": "Wulf A. Kaal",
  "identifier": "https://orcid.org/0009-0008-7840-1847"
 },
 "citation": "Wulf A. Kaal, How AI Models are Optimized Through Web3 Governance (2024). SSRN: https://ssrn.com/abstract=4855607",
 "isBasedOn": {
  "@type": "ScholarlyArticle",
  "name": "How AI Models are Optimized Through Web3 Governance",
  "datePublished": "2024",
  "url": "https://ssrn.com/abstract=4855607",
  "sha256": "eb0b3e62374b45a8fa888c6bde9725e606bcb46cf4b5e74a6e851d9f25099113",
  "encoding": "https://raw.githubusercontent.com/wulfkaal/Academic-Papers/main/papers/pdf/Kaal%20-%202024%20-%20How%20AI%20Models%20are%20Optimized%20Through%20Web3%20Governance.pdf"
 },
 "keywords": [
  "economics",
  "empirical-evidence"
 ],
 "about": [
  "reinforcement-learning",
  "sample-efficiency",
  "sparse-rewards",
  "sequence-generation",
  "human-learning"
 ],
 "abstract": "Deep RL methods often demand large amounts of training data, suggesting that the algorithms may differ fundamentally from those underlying human learning. Moreover, learning without supervision is particularly hard when the reward is sparse, which is likely to happen for sequence generation tasks.",
 "additionalProperty": [
  {
   "@type": "PropertyValue",
   "name": "claim_type",
   "value": "failure"
  },
  {
   "@type": "PropertyValue",
   "name": "confidence",
   "value": "evidenced"
  },
  {
   "@type": "PropertyValue",
   "name": "scope_conditions",
   "value": [
    "sparse reward settings",
    "sequence generation tasks"
   ]
  },
  {
   "@type": "PropertyValue",
   "name": "is_failure_mode",
   "value": true
  },
  {
   "@type": "PropertyValue",
   "name": "attestations",
   "value": "https://wulfkaal.github.io/colloquium/attestations/9eaf31289e4a67e6058e1a327b9f6cba08dabcc110d7f9faa6146393f3756bf2.json"
  },
  {
   "@type": "PropertyValue",
   "name": "edges",
   "value": []
  }
 ],
 "isPartOf": {
  "@id": "https://wulfkaal.github.io/claims/index.json"
 },
 "dateModified": "2026-07-27",
 "version": "1.0",
 "sha256": "9eaf31289e4a67e6058e1a327b9f6cba08dabcc110d7f9faa6146393f3756bf2",
 "canonicalForm": "https://wulfkaal.github.io/claims/4855607-014.md"
}