{
  "schemaVersion": 2,
  "fallacies": {
    "resourceId": "four-to-one",
    "revision": 1,
    "reviewed": "2026-09-25",
    "protocolVersion": 1,
    "articleSha256": "0b34f2e4e80c024b98aa76aa41290855caf93483ad27aded4080cb053cd2ef48",
    "status": "Completed for preserved text",
    "reviewer": "Codex-assisted E² editorial review; no independent second review",
    "scope": "All 35 retained sections, all 81 footnotes, and Appendices A–B. Repeated conclusions share one finding.",
    "exclusions": [
      "Images and embedded posts omitted from the edition",
      "Appendix C absent from the retrieved source",
      "Full independent verification of every cited work",
      "Independent linguistic verification of the Chinese translations"
    ],
    "rules": {
      "F-01": {
        "id": "F-01",
        "name": "Affirming the consequent",
        "criterion": "A deductive converse is asserted without support.",
        "version": 1,
        "source": {
          "url": "https://iep.utm.edu/fallacy/",
          "section": "Affirming the consequent",
          "accessed": "2026-09-25"
        }
      },
      "F-02": {
        "id": "F-02",
        "name": "Denying the antecedent",
        "criterion": "A deductive inverse is asserted without support.",
        "version": 1,
        "source": {
          "url": "https://iep.utm.edu/fallacy/",
          "section": "Denying the antecedent",
          "accessed": "2026-09-25"
        }
      },
      "I-01": {
        "id": "I-01",
        "name": "Equivocation",
        "criterion": "An inference depends on a term changing meaning.",
        "version": 1,
        "source": {
          "url": "https://iep.utm.edu/fallacy/",
          "section": "Equivocation",
          "accessed": "2026-09-25"
        }
      },
      "I-02": {
        "id": "I-02",
        "name": "False dilemma",
        "criterion": "Relevant alternatives are excluded without justification.",
        "version": 1,
        "source": {
          "url": "https://iep.utm.edu/fallacy/",
          "section": "False dilemma",
          "accessed": "2026-09-25"
        }
      },
      "I-03": {
        "id": "I-03",
        "name": "Circular reasoning",
        "criterion": "The disputed conclusion supplies its own support.",
        "version": 1,
        "source": {
          "url": "https://iep.utm.edu/fallacy/",
          "section": "Circular reasoning",
          "accessed": "2026-09-25"
        }
      },
      "I-04": {
        "id": "I-04",
        "name": "Straw man",
        "criterion": "A distorted attributed position is treated as refuted.",
        "version": 1,
        "source": {
          "url": "https://iep.utm.edu/fallacy/",
          "section": "Straw man",
          "accessed": "2026-09-25"
        }
      },
      "I-05": {
        "id": "I-05",
        "name": "Hasty generalization",
        "criterion": "A conclusion exceeds the demonstrated scope of its evidence.",
        "version": 1,
        "source": {
          "url": "https://iep.utm.edu/fallacy/",
          "section": "Hasty generalization",
          "accessed": "2026-09-25"
        }
      },
      "I-06": {
        "id": "I-06",
        "name": "False cause",
        "criterion": "A causal conclusion exceeds the evidence distinguishing alternatives.",
        "version": 1,
        "source": {
          "url": "https://iep.utm.edu/fallacy/",
          "section": "False cause",
          "accessed": "2026-09-25"
        }
      },
      "I-07": {
        "id": "I-07",
        "name": "Inappropriate authority",
        "criterion": "An appeal rests on irrelevant or unqualified authority.",
        "version": 1,
        "source": {
          "url": "https://iep.utm.edu/fallacy/",
          "section": "Appeal to Unqualified Authority",
          "accessed": "2026-09-25"
        }
      },
      "I-08": {
        "id": "I-08",
        "name": "Appeal to ignorance",
        "criterion": "Lack of evidence is treated as decisive proof.",
        "version": 1,
        "source": {
          "url": "https://iep.utm.edu/fallacy/",
          "section": "Appeal to ignorance",
          "accessed": "2026-09-25"
        }
      },
      "I-09": {
        "id": "I-09",
        "name": "Faulty comparison",
        "criterion": "A conclusion requires comparability that the quantities lack.",
        "version": 1,
        "source": {
          "url": "https://iep.utm.edu/fallacy/",
          "section": "Faulty comparison",
          "accessed": "2026-09-25"
        }
      },
      "I-10": {
        "id": "I-10",
        "name": "Accident",
        "criterion": "A qualified rule is applied without its necessary conditions.",
        "version": 1,
        "source": {
          "url": "https://iep.utm.edu/fallacy/",
          "section": "Accident",
          "accessed": "2026-09-25"
        }
      },
      "I-11": {
        "id": "I-11",
        "name": "Suppressed evidence",
        "criterion": "Material counterevidence is omitted from the inferential assessment.",
        "version": 1,
        "source": {
          "url": "https://iep.utm.edu/fallacy/",
          "section": "Suppressed evidence",
          "accessed": "2026-09-25"
        }
      },
      "I-12": {
        "id": "I-12",
        "name": "Ad hominem",
        "criterion": "A personal characterization substitutes for evaluating the argument.",
        "version": 1,
        "source": {
          "url": "https://iep.utm.edu/fallacy/",
          "section": "Ad hominem",
          "accessed": "2026-09-25"
        }
      },
      "I-13": {
        "id": "I-13",
        "name": "False analogy",
        "criterion": "An analogy supports a conclusion despite material disanalogies.",
        "version": 1,
        "source": {
          "url": "https://iep.utm.edu/fallacy/",
          "section": "False analogy",
          "accessed": "2026-09-25"
        }
      }
    },
    "findings": [
      {
        "id": "L01",
        "ruleId": "I-05",
        "status": "MATCH",
        "title": "Calibration across hardware",
        "section": "b-inference-arithmetic",
        "quote": "ratios between chips and models are sound",
        "claimIds": [
          "C16",
          "C18",
          "C19"
        ],
        "premises": [
          "The reported throughput bias is similar across the selected B200/B300 runs.",
          "Other deployment anchors produce a similar broad bias band."
        ],
        "conclusion": "Ratios across the modeled chips and models remain sound.",
        "inferentialMode": "Inductive",
        "implicitPremises": [],
        "predicates": [
          {
            "condition": "Evidence is taken from a limited set of systems",
            "value": true,
            "evidence": "Appendix B identifies 44 single-node V4 Pro FP4 runs on B200/B300."
          },
          {
            "condition": "Conclusion extends to untested comparisons",
            "value": true,
            "evidence": "The conclusion concerns ratios between chips and models, including the modeled Ascend systems."
          },
          {
            "condition": "No matching-error argument establishes this extension",
            "value": true,
            "evidence": "The text gives a common bias range, not matched residuals for each compared system."
          }
        ],
        "exclusion": {
          "condition": "The claim is limited to the tested pairings or supported by matched cross-system residuals",
          "value": false,
          "evidence": "The sentence states ratios between chips and models without that restriction."
        },
        "counterreading": "A planning heuristic could use common-mode error as an explicit assumption. The sentence presents ratio validity as a result, rather than an assumption.",
        "reason": "Similar errors on the tested systems do not establish cancellation on different architectures. Two biases within the same broad band can still distort their ratio.",
        "repair": "Restrict the conclusion to tested comparisons, or publish matched cross-system ratio residuals.",
        "revisionCondition": "Restrict the conclusion to tested comparisons, or publish matched cross-system ratio residuals.",
        "externalSources": [],
        "disposition": "Editorial finding",
        "reviewer": "Codex-assisted E² editorial review; no independent second review",
        "history": [
          {
            "date": "2026-09-25",
            "revision": 1,
            "decision": "MATCH"
          }
        ]
      },
      {
        "id": "L02",
        "ruleId": "I-09",
        "status": "MATCH",
        "title": "Capacity versus demand",
        "section": "inference-empiricals",
        "quote": "This is comfortably above the demand range, but not ridiculously so. Seems in the right ballpark.",
        "claimIds": [
          "C13",
          "C14",
          "C17",
          "C20"
        ],
        "premises": [
          "The Flash projection derates to an upper capacity near 23T total tokens/day.",
          "Router and revenue estimates cover a mixture of models and only part of total demand."
        ],
        "conclusion": "The fleet can comfortably serve the current demand.",
        "inferentialMode": "Inductive",
        "implicitPremises": [],
        "predicates": [
          {
            "condition": "The conclusion compares capacity with demand",
            "value": true,
            "evidence": "The section asks whether the fleet supports current demand."
          },
          {
            "condition": "Compared quantities represent different workloads or coverage",
            "value": true,
            "evidence": "Flash capacity is compared with mixed Flash/Pro demand; free traffic is explicitly outside the revenue estimate."
          },
          {
            "condition": "The mismatch affects the conclusion",
            "value": true,
            "evidence": "Total-token capacity changes with input mix, model allocation and service targets; uncovered demand can consume the margin."
          }
        ],
        "exclusion": {
          "condition": "A matched workload and complete demand bound justify the comparison",
          "value": false,
          "evidence": "The article gives neither a fleet allocation for the stated demand mix nor a bound on uncounted traffic."
        },
        "counterreading": "The phrase “ballpark” permits a rough plausibility check. It does not make the stronger fleet-sufficiency conclusion follow from unmatched quantities.",
        "reason": "The comparison does not establish fleet sufficiency: it uses an upper Flash capacity against partial mixed-model demand.",
        "repair": "Compare sustained capacity with a matched model/input mix, complete demand coverage and the same latency target.",
        "revisionCondition": "Compare sustained capacity with a matched model/input mix, complete demand coverage and the same latency target.",
        "externalSources": [],
        "disposition": "Editorial finding",
        "reviewer": "Codex-assisted E² editorial review; no independent second review",
        "history": [
          {
            "date": "2026-09-25",
            "revision": 1,
            "decision": "MATCH"
          }
        ]
      },
      {
        "id": "L03",
        "ruleId": "I-10",
        "status": "UNRESOLVED",
        "title": "Nine-month training ceiling",
        "section": "verdict",
        "quote": "which exceeds the theoretical 9-month Epoch ceiling",
        "claimIds": [
          "C21",
          "C22"
        ],
        "premises": [
          "Epoch estimates an economically preferred training duration under hardware and algorithmic progress.",
          "The proposed smaller fleet requires thirteen months."
        ],
        "conclusion": "A frontier-scale run on that fleet is less plausible.",
        "inferentialMode": "Inductive",
        "implicitPremises": [],
        "predicates": [
          {
            "condition": "A qualified rule is applied to a particular case",
            "value": true,
            "evidence": "The forecast applies Epoch’s duration result to a constrained hardware scenario."
          },
          {
            "condition": "Necessary conditions are discarded rather than used as a heuristic",
            "value": null,
            "evidence": "“Theoretical ceiling” sounds categorical; “less plausible” and the earlier economic explanation permit a qualified reading."
          }
        ],
        "exclusion": {
          "condition": "The duration is used only as a conditional planning benchmark",
          "value": null,
          "evidence": "Both the economic explanation and categorical ceiling wording appear in the article."
        },
        "counterreading": "A thirteen-month run can be less attractive without being impossible; the article explicitly calls the frontier-scale outcome less plausible.",
        "reason": "The ceiling wording overstates the source, but the surrounding argument can be read as a qualified economic comparison. That ambiguity prevents a confirmed fallacy label.",
        "repair": "Replace the universal ceiling language with the economic assumptions and test whether they hold for the lab.",
        "revisionCondition": "Replace the universal ceiling language with the economic assumptions and test whether they hold for the lab.",
        "externalSources": [
          {
            "url": "https://epoch.ai/data-insights/longest-training-run",
            "section": "Analysis",
            "accessed": "2026-09-25",
            "note": "The estimate uses opportunity costs from hardware and algorithmic progress."
          }
        ],
        "disposition": "Open interpretation",
        "reviewer": "Codex-assisted E² editorial review; no independent second review",
        "history": [
          {
            "date": "2026-09-25",
            "revision": 1,
            "decision": "UNRESOLVED"
          }
        ]
      },
      {
        "id": "L04",
        "ruleId": "I-01",
        "status": "UNRESOLVED",
        "title": "Generated versus total tokens",
        "section": "tl-dr",
        "quote": "~33 Flash (~8.6 Pro) Ttok/day of generated tokens",
        "claimIds": [
          "C13",
          "C14",
          "C20"
        ],
        "premises": [
          "The summary calls both output-only and input-plus-output quantities generated tokens.",
          "Later calculations distinguish total from output tokens."
        ],
        "conclusion": "The larger token figures support the fleet-capacity comparison.",
        "inferentialMode": "Inductive",
        "implicitPremises": [],
        "predicates": [
          {
            "condition": "A term has incompatible uses",
            "value": true,
            "evidence": "The summary uses generated for 1.9T output and roughly 33T total."
          },
          {
            "condition": "The inference depends on the changed meaning",
            "value": null,
            "evidence": "The later demand comparison uses total tokens and may not depend on this erroneous summary label."
          }
        ],
        "exclusion": {
          "condition": "This is a corrected labeling error rather than an inference using the ambiguity",
          "value": null,
          "evidence": "Inference: Theoreticals explicitly separates total and output tokens."
        },
        "counterreading": "The summary may contain a copy error. The later definitions let the reader recover the intended quantities.",
        "reason": "The labeling error is established in Claims. Its role in the reasoning is not clear enough to call it confirmed equivocation.",
        "repair": "Use total and output consistently, then reassess the actual capacity comparison.",
        "revisionCondition": "Use total and output consistently, then reassess the actual capacity comparison.",
        "externalSources": [],
        "disposition": "Open interpretation",
        "reviewer": "Codex-assisted E² editorial review; no independent second review",
        "history": [
          {
            "date": "2026-09-25",
            "revision": 1,
            "decision": "UNRESOLVED"
          }
        ]
      },
      {
        "id": "L05",
        "ruleId": "I-04",
        "status": "UNRESOLVED",
        "title": "The position attributed to critics",
        "section": "on-systems-thinking",
        "quote": "everything else is just commentary",
        "claimIds": [],
        "premises": [
          "The article describes a school of analysis that treats process nodes as decisive.",
          "Systems-level compensation can offset some chip-level disadvantages."
        ],
        "conclusion": "The described outlook is undermined by system-level compensation.",
        "inferentialMode": "Inductive",
        "implicitPremises": [],
        "predicates": [
          {
            "condition": "A position is attributed to an opposing group",
            "value": true,
            "evidence": "The paragraph defines the target school in strongly exclusive terms."
          },
          {
            "condition": "The attributed position materially distorts the identified opponent",
            "value": null,
            "evidence": "The general target is not delimited; cited commentary cannot establish what every intended target claims."
          }
        ],
        "exclusion": {
          "condition": "The target is an explicitly hypothetical extreme or a faithful sourced account",
          "value": null,
          "evidence": "The article may target only an extreme subset, but does not specify that subset precisely."
        },
        "counterreading": "Some critics may actually assert the extreme position. Others, including the cited Transformer argument, argue for delay rather than permanent prevention.",
        "reason": "There is a risk of refuting a stronger position than the cited delay argument. The broad target and omitted embedded posts prevent a definitive attribution finding.",
        "repair": "Name the exact proposition and source being rebutted, and test the strongest delay-and-cost version.",
        "revisionCondition": "Name the exact proposition and source being rebutted, and test the strongest delay-and-cost version.",
        "externalSources": [
          {
            "url": "https://www.transformernews.ai/p/deepseek-ceo-liang-wenfeng-export-controls-china",
            "section": "Opening argument",
            "accessed": "2026-09-25",
            "note": "The author advocates slowing development, not establishing permanent physical impossibility."
          }
        ],
        "disposition": "Open interpretation",
        "reviewer": "Codex-assisted E² editorial review; no independent second review",
        "history": [
          {
            "date": "2026-09-25",
            "revision": 1,
            "decision": "UNRESOLVED"
          }
        ]
      },
      {
        "id": "L06",
        "ruleId": "I-06",
        "status": "UNRESOLVED",
        "title": "Policy and economic damage",
        "section": "on-systems-thinking",
        "quote": "those negative externalities - like our present memory crunch - create objectively worse economic conditions at home and abroad",
        "claimIds": [],
        "premises": [
          "The article associates a policy outlook with constraints and shortages."
        ],
        "conclusion": "That outlook causes economic harm through the cited externalities.",
        "inferentialMode": "Inductive",
        "implicitPremises": [],
        "predicates": [
          {
            "condition": "A causal conclusion is offered",
            "value": true,
            "evidence": "The passage states that taking action on the outlook creates worse conditions."
          },
          {
            "condition": "A particular invalid causal inference is identifiable from the supplied text",
            "value": null,
            "evidence": "The causal chain is asserted, but the embedded supporting discussions are absent."
          }
        ],
        "exclusion": {
          "condition": "The cited supporting discussion supplies a discriminating causal mechanism or evidence",
          "value": null,
          "evidence": "The preserved text does not include the embedded posts."
        },
        "counterreading": "This may summarize an argument made in the missing supporting material, rather than infer causation from coincidence.",
        "reason": "The causal claim is under-supported in this edition. Missing premises alone do not establish a false-cause pattern.",
        "repair": "Supply the mechanism, comparison and relevant alternatives, then review the complete inference.",
        "revisionCondition": "Supply the mechanism, comparison and relevant alternatives, then review the complete inference.",
        "externalSources": [],
        "disposition": "Open interpretation",
        "reviewer": "Codex-assisted E² editorial review; no independent second review",
        "history": [
          {
            "date": "2026-09-25",
            "revision": 1,
            "decision": "UNRESOLVED"
          }
        ]
      }
    ],
    "coverage": [
      {
        "section": "tl-dr",
        "ruleIds": [
          "I-01",
          "I-05",
          "I-09"
        ],
        "outcome": "UNRESOLVED",
        "note": "Summary repeats L01/L02 and the L04 label ambiguity; the quarter-time discrepancy is arithmetic, not a separate fallacy.",
        "findingIds": [
          "L01",
          "L02",
          "L04"
        ]
      },
      {
        "section": "introduction",
        "ruleIds": [],
        "outcome": "NOT_APPLICABLE",
        "note": "Purpose and provenance commitments; no independent inference to classify.",
        "findingIds": []
      },
      {
        "section": "context",
        "ruleIds": [
          "I-02",
          "I-07",
          "I-08"
        ],
        "outcome": "NO_MATCH",
        "note": "Leak status and guessed topology are explicit. Questions list possibilities without asserting an exhaustive partition.",
        "findingIds": []
      },
      {
        "section": "on-systems-thinking",
        "ruleIds": [
          "I-04",
          "I-06",
          "I-12",
          "I-13"
        ],
        "outcome": "UNRESOLVED",
        "note": "L05/L06 remain open. The rail analogy illustrates a technical claim; hostile characterizations do not alone establish an ad hominem inference.",
        "findingIds": [
          "L05",
          "L06"
        ]
      },
      {
        "section": "enter-cortex-overclock",
        "ruleIds": [],
        "outcome": "NOT_APPLICABLE",
        "note": "Tool announcement and invitation.",
        "findingIds": []
      },
      {
        "section": "cortex",
        "ruleIds": [],
        "outcome": "NOT_APPLICABLE",
        "note": "Description of the author’s database and its labels.",
        "findingIds": []
      },
      {
        "section": "overclock",
        "ruleIds": [
          "I-05"
        ],
        "outcome": "NO_MATCH",
        "note": "Scaling is described as a modeling operation; empirical validity of every projection is not established by the description.",
        "findingIds": []
      },
      {
        "section": "methodology",
        "ruleIds": [
          "I-07"
        ],
        "outcome": "NO_MATCH",
        "note": "Technical conventions are attributed to relevant references; citation alone is not an inappropriate appeal.",
        "findingIds": []
      },
      {
        "section": "training",
        "ruleIds": [
          "F-01",
          "F-02",
          "I-03"
        ],
        "outcome": "NO_MATCH",
        "note": "A conditional dimensional calculation, not converse/inverse reasoning or proof by circular premise.",
        "findingIds": []
      },
      {
        "section": "inference",
        "ruleIds": [
          "F-01",
          "F-02",
          "I-03"
        ],
        "outcome": "NO_MATCH",
        "note": "An explicit approximate model. Its omissions require validation, not a fallacy label by themselves.",
        "findingIds": []
      },
      {
        "section": "batching",
        "ruleIds": [
          "I-11"
        ],
        "outcome": "NO_MATCH",
        "note": "The selection rule is stated; later disclosed latency error is a model-validation problem.",
        "findingIds": []
      },
      {
        "section": "calibration",
        "ruleIds": [
          "I-05",
          "I-11"
        ],
        "outcome": "UNRESOLVED",
        "note": "Observed errors are disclosed. The broad transfer claim is L01; without paired rows, selective reporting cannot be determined.",
        "findingIds": [
          "L01"
        ]
      },
      {
        "section": "assumptions",
        "ruleIds": [],
        "outcome": "NOT_APPLICABLE",
        "note": "The assumptions figure is omitted; retained text lists references. No missing figure content was inferred.",
        "findingIds": []
      },
      {
        "section": "4-1-spec-comparisons",
        "ruleIds": [],
        "outcome": "NOT_APPLICABLE",
        "note": "Section introduction.",
        "findingIds": []
      },
      {
        "section": "1-compute",
        "ruleIds": [
          "I-01",
          "I-04"
        ],
        "outcome": "UNRESOLVED",
        "note": "The author distinguishes numeric precision. Whether the criticism is represented fairly cannot be inferred from the chart-free excerpt alone; no additional finding asserted.",
        "findingIds": []
      },
      {
        "section": "2-memory",
        "ruleIds": [
          "I-01"
        ],
        "outcome": "NO_MATCH",
        "note": "GPU versus superchip units are explicitly distinguished.",
        "findingIds": []
      },
      {
        "section": "3-interconnect",
        "ruleIds": [
          "I-09"
        ],
        "outcome": "NO_MATCH",
        "note": "Scale and bandwidth are identified separately. Specifications require factual checks; no incompatible comparison is concealed in this list.",
        "findingIds": []
      },
      {
        "section": "4-system-vs-system",
        "ruleIds": [
          "I-09"
        ],
        "outcome": "NO_MATCH",
        "note": "Per-system and per-card comparisons are explicitly separated; unequal card counts are acknowledged.",
        "findingIds": []
      },
      {
        "section": "5-energy",
        "ruleIds": [
          "I-01",
          "I-09"
        ],
        "outcome": "NO_MATCH",
        "note": "System power and per-system ratios are scoped. Omitted chart numbers are not treated as inspected evidence.",
        "findingIds": []
      },
      {
        "section": "assessments",
        "ruleIds": [],
        "outcome": "NOT_APPLICABLE",
        "note": "Transition into scenarios.",
        "findingIds": []
      },
      {
        "section": "inference-theoreticals",
        "ruleIds": [
          "I-01",
          "I-11"
        ],
        "outcome": "NO_MATCH",
        "note": "Total/output quantities are separated here. SLA compliance remains unvalidated; that is not itself a demonstrated fallacy.",
        "findingIds": []
      },
      {
        "section": "inference-empiricals",
        "ruleIds": [
          "I-09",
          "I-05",
          "I-11"
        ],
        "outcome": "MATCH",
        "note": "L02 concerns workload and demand coverage. Claimed held-out calibration cannot be authenticated without the original pairs.",
        "findingIds": [
          "L02"
        ]
      },
      {
        "section": "post-training",
        "ruleIds": [
          "I-05"
        ],
        "outcome": "NO_MATCH",
        "note": "The text distinguishes rollout inference from gradient work. Fleet sufficiency inherits L02 rather than creating an independent finding.",
        "findingIds": [
          "L02"
        ]
      },
      {
        "section": "how-to-train-your-longcat-2-0",
        "ruleIds": [
          "F-01",
          "I-05",
          "I-07"
        ],
        "outcome": "NO_MATCH",
        "note": "A labeled hardware hypothesis is tested for feasibility, not deductively proven by matching duration. Vendor stability is qualified with “If true.”",
        "findingIds": []
      },
      {
        "section": "training-who-the-fuck-is-kimmy",
        "ruleIds": [
          "F-01",
          "I-05"
        ],
        "outcome": "NO_MATCH",
        "note": "The inference is probabilistic and cites independent cluster reporting; incomplete verification remains a provenance gap, not a formal converse error.",
        "findingIds": []
      },
      {
        "section": "training-a-future-deepseek-moment",
        "ruleIds": [
          "I-10"
        ],
        "outcome": "UNRESOLVED",
        "note": "Conditional scenario arithmetic is legitimate; applicability of the economic ceiling remains L03.",
        "findingIds": [
          "L03"
        ]
      },
      {
        "section": "can-china-train-frontier-models-on-domestic-chips",
        "ruleIds": [
          "I-01",
          "I-02"
        ],
        "outcome": "NO_MATCH",
        "note": "Alternative frontier definitions are explicitly distinguished; performance-per-dollar is an avowed preference.",
        "findingIds": []
      },
      {
        "section": "size-matters",
        "ruleIds": [
          "I-01"
        ],
        "outcome": "NO_MATCH",
        "note": "Parameter size, training compute and performance are kept separate in the text.",
        "findingIds": []
      },
      {
        "section": "performance-matters",
        "ruleIds": [
          "I-05",
          "I-08"
        ],
        "outcome": "NO_MATCH",
        "note": "The text separates post-training from pretraining and makes a possibility claim; it does not claim absence of a physical prohibition proves a completed run.",
        "findingIds": []
      },
      {
        "section": "verdict",
        "ruleIds": [
          "I-05",
          "I-10"
        ],
        "outcome": "UNRESOLVED",
        "note": "The forecast has explicit assumptions and uses “plausible.” Its nine-month premise retains the ambiguity in L03.",
        "findingIds": [
          "L03"
        ]
      },
      {
        "section": "closing-thoughts",
        "ruleIds": [
          "I-02",
          "I-06",
          "I-13"
        ],
        "outcome": "NO_MATCH",
        "note": "The ecosystems contrast is explicitly tentative and acknowledges tradeoffs. The rhetorical question does not assert a proven exhaustive partition; the NEV metaphor is not a demonstrated causal proof.",
        "findingIds": []
      },
      {
        "section": "footnotes",
        "ruleIds": [
          "I-01",
          "F-01"
        ],
        "outcome": "NO_MATCH",
        "note": "All 81 retained notes read. Footnote 55 explicitly limits an inference photo’s evidential role. Arithmetic and attribution issues remain in Claims; linked works were not all independently re-audited.",
        "findingIds": []
      },
      {
        "section": "appendix",
        "ruleIds": [],
        "outcome": "NOT_APPLICABLE",
        "note": "Heading for the two retained appendices.",
        "findingIds": []
      },
      {
        "section": "a-training-arithmetic-methodology",
        "ruleIds": [
          "I-03",
          "I-05"
        ],
        "outcome": "NO_MATCH",
        "note": "Conventions, assumptions, corrections and weak goodput sourcing are stated. Using an assumed efficiency in a scenario does not establish a sampled-population claim.",
        "findingIds": []
      },
      {
        "section": "b-inference-arithmetic",
        "ruleIds": [
          "I-05",
          "I-11"
        ],
        "outcome": "MATCH",
        "note": "L01 concerns cross-system ratio validity. The deployment model itself is conditional; incomplete benchmark records do not establish cherry-picking.",
        "findingIds": [
          "L01"
        ]
      }
    ],
    "screening": "All retained sections read for the initial ten patterns and five added patterns. Coverage records show the relevant candidate tests; this is not 35 × 15 independent tests.",
    "sources": [
      {
        "url": "https://iep.utm.edu/fallacy/",
        "title": "IEP · Fallacies",
        "accessed": "2026-09-25"
      },
      {
        "url": "https://plato.stanford.edu/entries/fallacies/",
        "title": "SEP · Fallacies",
        "accessed": "2026-09-25"
      }
    ]
  },
  "rubric": {
    "version": "1",
    "principles": [
      {
        "id": "define",
        "title": "Define the question"
      },
      {
        "id": "measure",
        "title": "Make the measure explicit"
      },
      {
        "id": "ask",
        "title": "Ask before asserting"
      },
      {
        "id": "separable",
        "title": "Keep the record separable"
      },
      {
        "id": "uncertainty",
        "title": "Preserve uncertainty"
      },
      {
        "id": "revision",
        "title": "Make revision possible"
      },
      {
        "id": "robust",
        "title": "Prefer robust action"
      }
    ],
    "bands": [
      [
        "A",
        "A clear, testable treatment"
      ],
      [
        "B",
        "A useful treatment with material gaps"
      ],
      [
        "C",
        "Significant gaps"
      ],
      [
        "D",
        "A weak treatment"
      ],
      [
        "F",
        "A failure to engage the principle"
      ]
    ]
  },
  "supportLevels": [
    [
      "Conjectural",
      "A possibility without direct supporting evidence."
    ],
    [
      "Reported",
      "An attributable assertion exists; the underlying event or quantity remains unverified."
    ],
    [
      "Supported",
      "Inspectable evidence supports the exact proposition, with material limits stated."
    ],
    [
      "Corroborated",
      "Independent evidence converges. Copies of one leak or repeated vendor claims do not count as independent sources."
    ],
    [
      "Established within scope",
      "Repeated checks or decisive evidence support a narrowly specified proposition under stated conditions. Still open to revision."
    ]
  ],
  "assessment": {
    "resourceId": "four-to-one",
    "rubricVersion": "1",
    "revision": 3,
    "reviewed": "2026-09-25",
    "history": [
      {
        "revision": 1,
        "date": "2026-09-25",
        "change": "Initial editorial assessment against the seven principles."
      },
      {
        "revision": 2,
        "date": "2026-09-25",
        "change": "Centralize the E² Grade record and document unperformed fallacy review. Grades unchanged."
      },
      {
        "revision": 3,
        "date": "2026-09-25",
        "change": "Review the full preserved text for fallacies, recording findings and unresolved readings. Migrate claim IDs to C01–C23. Grades unchanged."
      }
    ],
    "overall": "B",
    "summary": "The article exposes assumptions and model errors. Some conclusions exceed the evidence.",
    "counterargument": "The strongest competing explanation deserves a test: constraints can cause consequential delay without making eventual substitution impossible.",
    "principles": {
      "define": {
        "grade": "A−",
        "reason": "Separates compute, memory, systems, training and inference. Forecast criteria remain open."
      },
      "measure": {
        "grade": "B",
        "reason": "Careful denominators; inconsistent token labels and training-time headlines."
      },
      "ask": {
        "grade": "C+",
        "reason": "Tests technical alternatives. Gives less attention to the strongest opposing causal explanation."
      },
      "separable": {
        "grade": "B−",
        "reason": "Labels several assumptions. Sometimes turns feasibility into historical attribution."
      },
      "uncertainty": {
        "grade": "B+",
        "reason": "Discloses large residuals. Conclusions do not always carry them through."
      },
      "revision": {
        "grade": "B+",
        "reason": "Records corrections. The exact comparison dataset and model version remain unavailable here."
      },
      "robust": {
        "grade": "C+",
        "reason": "Identifies multiple constraints; does not compare concrete decisions across scenarios. Partly outside the article’s scope."
      }
    },
    "fallacyReview": {
      "status": "Completed for preserved text",
      "record": "four-to-one-fallacies.json",
      "scope": "Text, footnotes and Appendices A–B reviewed. Missing media and Appendix C excluded; unresolved readings retained."
    }
  },
  "review": {
    "resourceId": "four-to-one",
    "revision": 3,
    "reviewed": "2026-09-25",
    "weightPolicy": "Equal weight per ledger claim (1). No importance weighting has been assigned.",
    "provenance": [
      {
        "id": "attributed",
        "label": "Article attribution",
        "description": "The claim is traceable to the article. Its exact proposition is not independently established here.",
        "claims": [
          "C01",
          "C11",
          "C15",
          "C16",
          "C18",
          "C19",
          "C20",
          "C21",
          "C22"
        ]
      },
      {
        "id": "primary",
        "label": "First-party source checked",
        "description": "The exact report is traceable to an inspected vendor, model-card or research source. This verifies the attribution, not the underlying event independently.",
        "claims": [
          "C03",
          "C04",
          "C09",
          "C10"
        ]
      },
      {
        "id": "conditional",
        "label": "Conditional calculation",
        "description": "Inputs and arithmetic are identified in the article. Availability and outcome of reproduction are counted separately below.",
        "claims": [
          "C05",
          "C06",
          "C07",
          "C12",
          "C13",
          "C14",
          "C17"
        ]
      },
      {
        "id": "assumed",
        "label": "Declared assumption",
        "description": "A chosen input, not an observation.",
        "claims": [
          "C02",
          "C08"
        ]
      },
      {
        "id": "value",
        "label": "Declared value judgment",
        "description": "An objective or preference, outside the empirical support ladder.",
        "claims": [
          "C23"
        ]
      }
    ],
    "weights": {
      "C01": 1,
      "C02": 1,
      "C03": 1,
      "C04": 1,
      "C05": 1,
      "C06": 1,
      "C07": 1,
      "C08": 1,
      "C09": 1,
      "C10": 1,
      "C11": 1,
      "C12": 1,
      "C13": 1,
      "C14": 1,
      "C15": 1,
      "C16": 1,
      "C17": 1,
      "C18": 1,
      "C19": 1,
      "C20": 1,
      "C21": 1,
      "C22": 1,
      "C23": 1
    },
    "checks": [
      {
        "claimId": "C05",
        "kind": "arithmetic",
        "outcome": "matched",
        "note": "6N training FLOPs match the printed inputs.",
        "resultPaths": [
          "longcat.flops"
        ]
      },
      {
        "claimId": "C06",
        "kind": "arithmetic",
        "outcome": "matched",
        "note": "Base duration and accelerator-days reproduced.",
        "resultPaths": [
          "longcat.baseDays",
          "longcat.acceleratorDays"
        ]
      },
      {
        "claimId": "C07",
        "kind": "arithmetic",
        "outcome": "discrepancy",
        "note": "The caption and appendix vary different inputs.",
        "resultPaths": [
          "longcat.fixedMfuRangeDays",
          "longcat.jointMfuGoodputRangeDays"
        ]
      },
      {
        "claimId": "C12",
        "kind": "arithmetic",
        "outcome": "discrepancy",
        "note": "53/99 does not equal one quarter.",
        "resultPaths": [
          "hypotheticalDeepSeek.timeRatio"
        ]
      },
      {
        "claimId": "C13",
        "kind": "arithmetic",
        "outcome": "discrepancy",
        "note": "The larger numbers include input tokens.",
        "resultPaths": [
          "tokens.flashTotal16to1",
          "tokens.proTotal16to1"
        ]
      },
      {
        "claimId": "C14",
        "kind": "arithmetic",
        "outcome": "discrepancy",
        "note": "The two totals use different workload mixes.",
        "resultPaths": [
          "tokens.proTotal16to1",
          "tokens.proTotal26_8to1"
        ]
      },
      {
        "claimId": "C15",
        "kind": "model-comparison",
        "outcome": "unavailable",
        "note": "Executable comparison cases are missing.",
        "resultPaths": []
      },
      {
        "claimId": "C16",
        "kind": "empirical-comparison",
        "outcome": "unavailable",
        "note": "The exact cohort and paired predictions are missing.",
        "resultPaths": []
      },
      {
        "claimId": "C17",
        "kind": "service-validation",
        "outcome": "unvalidated",
        "note": "Latency and batching require a new validation run. Sensitivity arithmetic does not validate the service target.",
        "resultPaths": []
      },
      {
        "claimId": "C18",
        "kind": "empirical-comparison",
        "outcome": "unavailable",
        "note": "Paired interactivity measurements and predictions are missing.",
        "resultPaths": []
      }
    ]
  },
  "claims": [
    {
      "id": "C01",
      "text": "DeepSeek received about 16,000 Ascend 950 chips.",
      "type": "Alleged leak",
      "support": "Reported",
      "review": "Not independently verified",
      "section": "context",
      "sourceIds": [
        "article"
      ],
      "dependencies": [],
      "reason": "The article attributes this allocation to reporting about leaked remarks. Neither the original meeting record nor delivery records were authenticated in this review. Repetition of the report is not independent corroboration.",
      "revise": "An authenticated primary account or delivery record identifying date, quantity and chip variant."
    },
    {
      "id": "C02",
      "text": "The allocation is 950DT hardware in two Atlas 950 SuperPoDs.",
      "type": "Model assumption",
      "support": null,
      "review": "Source checked",
      "section": "context",
      "sourceIds": [
        "article",
        "huawei",
        "atlas"
      ],
      "dependencies": [
        "C01",
        "C03"
      ],
      "reason": "The author marks the topology as a guess. Huawei’s roadmap schedules 950DT and Atlas 950 for Q4 2026, after the article’s August publication. An early allocation is possible; advertised maximum size does not establish the delivered SKU or topology.",
      "revise": "A deployment disclosure specifying SKU, system layout and operational date."
    },
    {
      "id": "C03",
      "text": "Huawei advertised an Atlas 950 system with up to 8,192 NPUs.",
      "type": "Reported fact",
      "support": "Supported",
      "review": "Source checked",
      "section": "3-interconnect",
      "sourceIds": [
        "atlas"
      ],
      "dependencies": [],
      "reason": "The vendor announcement supports this narrow statement about advertised scale. It does not measure workload throughput or verify DeepSeek’s allocation.",
      "revise": "A revised vendor specification. Actual performance needs separate workload measurements."
    },
    {
      "id": "C04",
      "text": "The LongCat model card reports about 48B active parameters and more than 35T training tokens.",
      "type": "Reported fact",
      "support": "Supported",
      "review": "Source checked",
      "section": "how-to-train-your-longcat-2-0",
      "sourceIds": [
        "longcat"
      ],
      "dependencies": [],
      "reason": "The first-party card supports the reported architecture and corpus scale. The 35T scenario uses a rounded lower input, not the exact disclosed training total. The underlying training logs were not independently measured.",
      "revise": "Versioned architecture details or logs that revise the active count or exact token total."
    },
    {
      "id": "C05",
      "text": "At 48B active parameters and 35T tokens, the 6N training estimate is 1.008 × 10²⁵ FLOPs.",
      "type": "Calculated output",
      "support": null,
      "review": "Arithmetic reproduced",
      "section": "a-training-arithmetic-methodology",
      "sourceIds": [
        "article",
        "scaling"
      ],
      "dependencies": [
        "C04"
      ],
      "reason": "6 × 48 × 10⁹ × 35 × 10¹² = 1.008 × 10²⁵. This is a conditional estimate of model FLOPs, not measured device work. Efficiency must use the same FLOP convention to avoid double-counting overhead.",
      "revise": "Changed inputs, or an explicit architecture-specific accounting with a consistently defined efficiency denominator."
    },
    {
      "id": "C06",
      "text": "The printed LongCat base scenario takes 32.41 days on 50,000 chips.",
      "type": "Calculated output",
      "support": null,
      "review": "Arithmetic reproduced",
      "section": "a-training-arithmetic-methodology",
      "sourceIds": [
        "article"
      ],
      "dependencies": [
        "C05",
        "C08"
      ],
      "reason": "At 400 trillion FLOP/s per chip, 30% MFU and 60% goodput, the result is 1.620 million accelerator-days. This does not identify the historical chip SKU or actual duration. A numerical TFLOP/s input requires a factor of 10¹² in the denominator.",
      "revise": "A different dense peak, efficiency denominator, training volume or measured run duration."
    },
    {
      "id": "C07",
      "text": "The LongCat caption’s 22–98-day range follows from fixed 30% MFU and 40–75% goodput.",
      "type": "Calculated output",
      "support": null,
      "review": "Unresolved conflict",
      "section": "how-to-train-your-longcat-2-0",
      "sourceIds": [
        "article"
      ],
      "dependencies": [
        "C05"
      ],
      "reason": "Those fixed-MFU inputs yield 25.93–48.61 days. The wider 22.22–97.22-day range requires varying MFU from 15% to 35% as well. The caption and Appendix A describe different sweeps; neither range is a confidence interval.",
      "revise": "Correct the caption to name both varying inputs, or replace its range with the fixed-MFU result."
    },
    {
      "id": "C08",
      "text": "Ascend goodput is 40–75% across the modeled training scenarios.",
      "type": "Model assumption",
      "support": null,
      "review": "Source checked",
      "section": "a-training-arithmetic-methodology",
      "sourceIds": [
        "article"
      ],
      "dependencies": [],
      "reason": "The author explicitly calls this weakly sourced. It is a sensitivity range, not a measured fleet statistic. Holding it and MFU constant when scaling from thousands to hundreds of thousands of chips is a further assumption.",
      "revise": "Measured productive wall-clock fractions by fleet size, topology, workload and restart policy."
    },
    {
      "id": "C09",
      "text": "Pangu Ultra MoE’s authors report 30.0% MFU on 6,000 Ascend NPUs.",
      "type": "Reported measurement",
      "support": "Supported",
      "review": "Source checked",
      "section": "a-training-arithmetic-methodology",
      "sourceIds": [
        "pangu"
      ],
      "dependencies": [],
      "reason": "The abstract supports this first-party measurement report. It anchors one workload and scale; it does not independently validate the article’s broader Ascend fleet scenarios.",
      "revise": "Independent reproduction, or detailed counterevidence about the workload and MFU denominator."
    },
    {
      "id": "C10",
      "text": "Kimi K3’s model card reports 104B active parameters.",
      "type": "Reported fact",
      "support": "Supported",
      "review": "Source checked",
      "section": "training-who-the-fuck-is-kimmy",
      "sourceIds": [
        "kimi"
      ],
      "dependencies": [],
      "reason": "The model summary reports this architecture figure. Combined with an assumed 35T tokens, the 6N calculation yields 2.184 × 10²⁵ FLOPs. The assumed corpus does not become a reported training fact.",
      "revise": "A revised model card or a disclosed training-token total."
    },
    {
      "id": "C11",
      "text": "Kimi K3 was most likely trained on foreign silicon.",
      "type": "Inference",
      "support": "Reported",
      "review": "Not independently verified",
      "section": "training-who-the-fuck-is-kimmy",
      "sourceIds": [
        "article"
      ],
      "dependencies": [
        "C10"
      ],
      "reason": "The article cites reporting about a Hopper cluster. A plausible 47-day modeled schedule alone cannot establish which machines trained the model. The cited historical reporting was not authenticated in this audit.",
      "revise": "Primary training provenance, allocation records or a technical report naming the fleet."
    },
    {
      "id": "C12",
      "text": "50,000 GB300s take one quarter of the cited 99-day Ascend run.",
      "type": "Calculated output",
      "support": null,
      "review": "Unresolved conflict",
      "section": "training-a-future-deepseek-moment",
      "sourceIds": [
        "article"
      ],
      "dependencies": [],
      "reason": "The body gives 53 days, so its time ratio is 53/99 = 0.535. One quarter would be 24.75 days. Confirm model version, precision, efficiency and GPU-versus-superchip units before choosing a replacement headline.",
      "revise": "A consistent configuration export that reconciles the summary and body."
    },
    {
      "id": "C13",
      "text": "The 33T Flash / 8.6T Pro headline denotes generated tokens per day.",
      "type": "Calculated output",
      "support": null,
      "review": "Contradicted as labeled",
      "section": "tl-dr",
      "sourceIds": [
        "article"
      ],
      "dependencies": [],
      "reason": "At the printed 16:1 input:output mix, 1.9T Flash output implies 32.3T total; 0.5T Pro output implies 8.5T total. The larger numbers include input tokens. “Generated” is the wrong label under those inputs.",
      "revise": "Relabel total tokens and keep generated/output tokens separate throughout."
    },
    {
      "id": "C14",
      "text": "The Pro body’s roughly 14T total is the same 16:1 scenario as the opening.",
      "type": "Calculated output",
      "support": null,
      "review": "Unresolved conflict",
      "section": "inference-theoreticals",
      "sourceIds": [
        "article"
      ],
      "dependencies": [
        "C13"
      ],
      "reason": "At 0.5T output, 14T total instead matches the 26.8:1 Pro mix in footnote 39: 0.5 × 27.8 = 13.9. Changing input mix also changes prefill work; multiplying a fixed output projection is only a reconciliation of labels, not a rerun.",
      "revise": "Name the workload mix for each result and recompute sustained throughput at that mix."
    },
    {
      "id": "C15",
      "text": "Overclock reproduces Tensor Economics calculations within 2–5%.",
      "type": "Calculated comparison",
      "support": "Reported",
      "review": "Not reproduced",
      "section": "calibration",
      "sourceIds": [
        "article"
      ],
      "dependencies": [],
      "reason": "This is author-reported implementation agreement. Even if reproduced, agreement with another model is not independent measurement of deployment performance.",
      "revise": "Publish reference cases, inputs, expected values and executable comparison code."
    },
    {
      "id": "C16",
      "text": "Across 44 InferenceX runs, Overclock overpredicts throughput by median 1.81×, with p10–p90 of 1.39–2.18×.",
      "type": "Comparison with independent measurements",
      "support": "Reported",
      "review": "Not reproduced",
      "section": "inference-empiricals",
      "sourceIds": [
        "article",
        "infx",
        "infxCode"
      ],
      "dependencies": [],
      "reason": "The author says these measurements were held out. The public historical query is retained with this audit, but the exact 44-row selection and paired Overclock predictions are unavailable here. We cannot recompute the median, quantiles or flatness by concurrency from aggregate statements.",
      "revise": "Publish result IDs, paired predictions, metric units, versions, exclusions and the tuning/holdout split."
    },
    {
      "id": "C17",
      "text": "The stated fleet capacities meet 20 output tokens per second per user.",
      "type": "Calculated output",
      "support": null,
      "review": "Not validated",
      "section": "inference-theoreticals",
      "sourceIds": [
        "article"
      ],
      "dependencies": [
        "C02",
        "C16",
        "C18"
      ],
      "reason": "The model selects its batch using predicted interactivity. The article also reports a 2.8× median overestimate of that metric. Dividing a projected 20 by 2.8 gives 7.14 as an illustrative sensitivity, not a pointwise correction. Aggregate throughput derating cannot validate the batch-selection constraint.",
      "revise": "Recalibrate latency, rerun the batch sweep, and validate the selected configurations against measured service targets."
    },
    {
      "id": "C18",
      "text": "Overclock’s per-user interactivity is 2.8× high at the median.",
      "type": "Comparison with independent measurements",
      "support": "Reported",
      "review": "Not reproduced",
      "section": "calibration",
      "sourceIds": [
        "article"
      ],
      "dependencies": [
        "C16"
      ],
      "reason": "Disclosing this residual is a strength. The figure still requires paired data and an explicit latency statistic. Its effect on the selected operating point matters more than whether the aggregate-throughput error looks stable.",
      "revise": "Publish the per-row interactivity residuals and evaluate held-out configurations after recalibration."
    },
    {
      "id": "C19",
      "text": "Cross-chip and cross-model ratios remain sound despite absolute throughput errors.",
      "type": "Inference",
      "support": "Conjectural",
      "review": "Not established",
      "section": "b-inference-arithmetic",
      "sourceIds": [
        "article"
      ],
      "dependencies": [
        "C16"
      ],
      "reason": "Stable bias on B200/B300 does not establish matched bias on Ascend superpods or other models. If two biases independently occupy 1.4–2.2, a predicted ratio can be 0.64–1.57 times the true ratio. This is a sensitivity bound, not an empirical interval.",
      "revise": "Matched measurements across the compared systems, workload mixes and precisions, with ratio residuals reported."
    },
    {
      "id": "C20",
      "text": "The 16,000-chip fleet comfortably covers total DeepSeek demand.",
      "type": "Inference",
      "support": "Conjectural",
      "review": "Not established",
      "section": "inference-empiricals",
      "sourceIds": [
        "article"
      ],
      "dependencies": [
        "C01",
        "C02",
        "C13",
        "C14",
        "C17"
      ],
      "reason": "Flash capacity is compared with mixed-model demand. Router traffic is a subset; revenue-based estimates omit free usage. Price, cache, workload mix and peak load also matter. The printed figures do not establish enough capacity for total demand at the service target.",
      "revise": "A matched output/input workload mix, full demand coverage, measured sustainable capacity and peak-load margin."
    },
    {
      "id": "C21",
      "text": "Nine months is a ceiling beyond which a training run is not viable.",
      "type": "Inference",
      "support": "Reported",
      "review": "Unresolved interpretation",
      "section": "training-a-future-deepseek-moment",
      "sourceIds": [
        "article",
        "epoch"
      ],
      "dependencies": [],
      "reason": "Epoch estimates an economic optimum near 8.6 months with a 6.1–14.2-month interval under its assumptions. It is not a physical ceiling. Restricted access to future hardware can change the opportunity cost of waiting.",
      "revise": "An economic calculation using the actual lab’s hardware access, algorithmic progress and release incentives."
    },
    {
      "id": "C22",
      "text": "A predominantly domestic frontier-performance run is plausible within twelve months of publication.",
      "type": "Forecast",
      "support": "Conjectural",
      "review": "Not yet resolvable",
      "section": "verdict",
      "sourceIds": [
        "article"
      ],
      "dependencies": [
        "C08",
        "C17",
        "C19"
      ],
      "reason": "The implied deadline is August 6, 2027. The article gives no fixed benchmark suite, comparator set, domestic-content threshold or qualifying training stage. Without these, later events can be fitted to the forecast after the fact.",
      "revise": "Preregister resolution criteria and a probability, then score against the August 6, 2027 outcome."
    },
    {
      "id": "C23",
      "text": "Performance per dollar matters most when defining the frontier.",
      "type": "Value judgment",
      "support": null,
      "review": "Scope identified",
      "section": "can-china-train-frontier-models-on-domestic-chips",
      "sourceIds": [
        "article"
      ],
      "dependencies": [],
      "reason": "This selects an objective. It can be appropriate for a user or operator, but it is not a universal empirical result. Capability, latency, reliability and affordability can lead different users to different choices.",
      "revise": "State the decision-maker and their objective; compare decisions across plausible priorities."
    }
  ],
  "sources": {
    "article": {
      "title": "Original article",
      "url": "https://x.com/rydcunningham/status/2085434224178303376",
      "location": "Body and Appendices A–B; retrieved September 25, 2026"
    },
    "huawei": {
      "title": "Huawei Connect 2025 keynote",
      "url": "https://www.huawei.com/en/news/2025/9/hc-xu-keynote-speech",
      "location": "Ascend 950 / 960 roadmap"
    },
    "atlas": {
      "title": "Huawei MWC 2026 announcement",
      "url": "https://www.huawei.com/en/news/2026/3/mwc-superpod-ai",
      "location": "Atlas 950 SuperPoD specifications"
    },
    "longcat": {
      "title": "LongCat 2.0 model card",
      "url": "https://huggingface.co/meituan-longcat/LongCat-2.0",
      "location": "Introduction; card retrieved September 25, 2026"
    },
    "kimi": {
      "title": "Kimi K3 model card",
      "url": "https://huggingface.co/moonshotai/Kimi-K3",
      "location": "Model summary; card retrieved September 25, 2026"
    },
    "pangu": {
      "title": "Pangu Ultra MoE report",
      "url": "https://arxiv.org/abs/2505.04519",
      "location": "Abstract: 30.0% MFU on 6K Ascend NPUs"
    },
    "infx": {
      "title": "InferenceX API documentation",
      "url": "https://inferencex.semianalysis.com/api",
      "location": "Historical inference view, metric and configuration definitions"
    },
    "infxCode": {
      "title": "Official InferenceX repository",
      "url": "https://github.com/SemiAnalysisAI/InferenceX",
      "location": "Benchmark implementation and reproducibility materials"
    },
    "epoch": {
      "title": "Epoch: longest training runs",
      "url": "https://epoch.ai/data-insights/longest-training-run",
      "location": "Estimated economically optimal training duration"
    },
    "scaling": {
      "title": "How to Scale Your Model",
      "url": "https://jax-ml.github.io/scaling-book/transformers/",
      "location": "Transformer training FLOP approximation"
    }
  },
  "audit": {
    "reviewed": "2026-09-25",
    "scope": "Independent arithmetic and sensitivity checks of printed inputs; not a replication of Overclock or its 44 paired benchmark comparisons.",
    "longcat": {
      "inputs": {
        "active": 48000000000,
        "tokens": 35000000000000,
        "cards": 50000,
        "denseFlops": 400000000000000,
        "mfu": 0.3,
        "goodput": 0.6
      },
      "flops": 1.008e+25,
      "baseDays": 32.40740740740741,
      "acceleratorDays": 1620370.3703703706,
      "jointMfuGoodputRangeDays": [
        22.22222222222222,
        97.22222222222223
      ],
      "fixedMfuRangeDays": [
        25.925925925925927,
        48.611111111111114
      ]
    },
    "kimi": {
      "assumedTokens": 35000000000000,
      "active": 104000000000,
      "flops": 2.184e+25
    },
    "hypotheticalDeepSeek": {
      "active": 800000000000,
      "assumedTokens": 32000000000000,
      "flops": 1.536e+26,
      "printedGb300Days": 53,
      "printedAscendDays": 99,
      "timeRatio": 0.5353535353535354,
      "impliedPerCardSpeedRatio": 7.471698113207547
    },
    "tokens": {
      "unit": "trillion tokens/day",
      "flashOutput": 1.9,
      "flashTotal16to1": 32.3,
      "proOutput": 0.5,
      "proTotal16to1": 8.5,
      "proTotal26_8to1": 13.9
    },
    "calibrationSensitivity": {
      "note": "Illustrations conditional on reported errors; not fitted corrections or confidence intervals. Rerun batching and prefill before using operationally.",
      "flashTotalDerated": [
        14.68181818181818,
        23.07142857142857
      ],
      "proTotalDerated": [
        6.3181818181818175,
        9.928571428571429
      ],
      "projected20TokensPerSecondDividedByReportedMedianBias": 7.142857142857143,
      "crossSystemRatioErrorIfBiasesVaryIndependently": [
        0.6363636363636362,
        1.5714285714285716
      ]
    },
    "calibrationReplication": {
      "status": "not reproduced",
      "missing": [
        "Exact 44 benchmark result IDs and snapshot",
        "Paired Overclock predictions and versioned configuration",
        "Metric definitions, GPU/group units, latency statistic, and tuning/holdout split"
      ]
    }
  },
  "statistics": {
    "total": 23,
    "empirical": 13,
    "excluded": 10,
    "totalWeight": 13,
    "weightedSum": 14,
    "weightedSupport": 1.0769230769230769,
    "support": [
      {
        "label": "Reported",
        "score": 1,
        "count": 6
      },
      {
        "label": "Supported",
        "score": 2,
        "count": 4
      },
      {
        "label": "Conjectural",
        "score": 0,
        "count": 3
      },
      {
        "label": "Corroborated",
        "score": 3,
        "count": 0
      },
      {
        "label": "Established within scope",
        "score": 4,
        "count": 0
      }
    ],
    "types": [
      {
        "label": "Calculated output",
        "count": 7
      },
      {
        "label": "Inference",
        "count": 4
      },
      {
        "label": "Reported fact",
        "count": 3
      },
      {
        "label": "Model assumption",
        "count": 2
      },
      {
        "label": "Comparison with independent measurements",
        "count": 2
      },
      {
        "label": "Alleged leak",
        "count": 1
      },
      {
        "label": "Reported measurement",
        "count": 1
      },
      {
        "label": "Calculated comparison",
        "count": 1
      },
      {
        "label": "Forecast",
        "count": 1
      },
      {
        "label": "Value judgment",
        "count": 1
      }
    ],
    "reviews": [
      {
        "label": "Source checked",
        "count": 6
      },
      {
        "label": "Unresolved conflict",
        "count": 3
      },
      {
        "label": "Not reproduced",
        "count": 3
      },
      {
        "label": "Not independently verified",
        "count": 2
      },
      {
        "label": "Arithmetic reproduced",
        "count": 2
      },
      {
        "label": "Not established",
        "count": 2
      },
      {
        "label": "Contradicted as labeled",
        "count": 1
      },
      {
        "label": "Not validated",
        "count": 1
      },
      {
        "label": "Unresolved interpretation",
        "count": 1
      },
      {
        "label": "Not yet resolvable",
        "count": 1
      },
      {
        "label": "Scope identified",
        "count": 1
      }
    ],
    "reviewGroups": [
      {
        "id": "unverified",
        "label": "Unverified",
        "description": "Further evidence, clarification or validation is needed. Each claim records the required check.",
        "count": 9,
        "claims": [
          "C01",
          "C11",
          "C15",
          "C16",
          "C17",
          "C18",
          "C19",
          "C20",
          "C21"
        ],
        "states": [
          {
            "label": "Not reproduced",
            "count": 3,
            "claims": [
              "C15",
              "C16",
              "C18"
            ]
          },
          {
            "label": "Not independently verified",
            "count": 2,
            "claims": [
              "C01",
              "C11"
            ]
          },
          {
            "label": "Not established",
            "count": 2,
            "claims": [
              "C19",
              "C20"
            ]
          },
          {
            "label": "Not validated",
            "count": 1,
            "claims": [
              "C17"
            ]
          },
          {
            "label": "Unresolved interpretation",
            "count": 1,
            "claims": [
              "C21"
            ]
          }
        ]
      },
      {
        "id": "verified",
        "label": "Verified within scope",
        "description": "The attribution or conditional arithmetic checks out; underlying events and inputs may remain unverified.",
        "count": 6,
        "claims": [
          "C03",
          "C04",
          "C05",
          "C06",
          "C09",
          "C10"
        ],
        "states": [
          {
            "label": "Source checked",
            "count": 4,
            "claims": [
              "C03",
              "C04",
              "C09",
              "C10"
            ]
          },
          {
            "label": "Arithmetic reproduced",
            "count": 2,
            "claims": [
              "C05",
              "C06"
            ]
          }
        ]
      },
      {
        "id": "disputed",
        "label": "Disputed",
        "description": "A contradiction or unresolved conflict was found.",
        "count": 4,
        "claims": [
          "C07",
          "C12",
          "C13",
          "C14"
        ],
        "states": [
          {
            "label": "Unresolved conflict",
            "count": 3,
            "claims": [
              "C07",
              "C12",
              "C14"
            ]
          },
          {
            "label": "Contradicted as labeled",
            "count": 1,
            "claims": [
              "C13"
            ]
          }
        ]
      },
      {
        "id": "not-applicable",
        "label": "Not applicable",
        "description": "Chosen assumptions or preferences, rather than factual findings.",
        "count": 3,
        "claims": [
          "C02",
          "C08",
          "C23"
        ],
        "states": [
          {
            "label": "Source checked",
            "count": 2,
            "claims": [
              "C02",
              "C08"
            ]
          },
          {
            "label": "Scope identified",
            "count": 1,
            "claims": [
              "C23"
            ]
          }
        ]
      },
      {
        "id": "pending",
        "label": "Pending",
        "description": "A future outcome; verification requires a deadline and fixed resolution criteria.",
        "count": 1,
        "claims": [
          "C22"
        ],
        "states": [
          {
            "label": "Not yet resolvable",
            "count": 1,
            "claims": [
              "C22"
            ]
          }
        ]
      }
    ],
    "provenance": [
      {
        "id": "attributed",
        "label": "Article attribution",
        "description": "The claim is traceable to the article. Its exact proposition is not independently established here.",
        "claims": [
          "C01",
          "C11",
          "C15",
          "C16",
          "C18",
          "C19",
          "C20",
          "C21",
          "C22"
        ],
        "count": 9
      },
      {
        "id": "conditional",
        "label": "Conditional calculation",
        "description": "Inputs and arithmetic are identified in the article. Availability and outcome of reproduction are counted separately below.",
        "claims": [
          "C05",
          "C06",
          "C07",
          "C12",
          "C13",
          "C14",
          "C17"
        ],
        "count": 7
      },
      {
        "id": "primary",
        "label": "First-party source checked",
        "description": "The exact report is traceable to an inspected vendor, model-card or research source. This verifies the attribution, not the underlying event independently.",
        "claims": [
          "C03",
          "C04",
          "C09",
          "C10"
        ],
        "count": 4
      },
      {
        "id": "assumed",
        "label": "Declared assumption",
        "description": "A chosen input, not an observation.",
        "claims": [
          "C02",
          "C08"
        ],
        "count": 2
      },
      {
        "id": "value",
        "label": "Declared value judgment",
        "description": "An objective or preference, outside the empirical support ladder.",
        "claims": [
          "C23"
        ],
        "count": 1
      }
    ],
    "quantitative": 10,
    "auditOutcomes": {
      "matched": 2,
      "discrepancy": 4,
      "unavailable": 3,
      "unvalidated": 1
    },
    "empiricalChecks": 2,
    "empiricalReplicated": 0,
    "dependencies": 21
  }
}
