{
  "$schema": "https://themachinepress.com/schemas/edition-v1.schema.json",
  "schema_version": "1.0.0",
  "document_type": "machine_press_edition",
  "edition_id": "mp-2026-09-06-morning-0059",
  "source_edition_id": "mp-2026-09-06-morning-0059",
  "edition_number": 59,
  "edition_label": "Morning edition",
  "status": "published",
  "revision": 1,
  "published_at": "2026-09-06T09:00:00.000-04:00",
  "modified_at": "2026-09-06T09:00:00.000-04:00",
  "timezone": "America/Detroit",
  "canonical_url": "https://themachinepress.com/edition/2026-09-06",
  "html_url": "https://themachinepress.com/edition/2026-09-06",
  "json_url": "https://themachinepress.com/edition/2026-09-06.json",
  "markdown_url": "https://themachinepress.com/edition/2026-09-06.md",
  "lead_story_id": "mp-2026-09-06-001",
  "lead_summary": "Rollout-based advantage tests found that language-model judges could identify consequential chain-of-thought steps better than chance—but far below the experiment's noise ceiling.",
  "coverage_window": {
    "start": "2026-09-03T03:52:56.000-04:00",
    "end": "2026-09-06T08:20:00.000-04:00"
  },
  "story_count": 27,
  "stories": [
    {
      "story_id": "mp-2026-09-06-001",
      "source_story_id": "tmp-lead-reasoning-step-importance",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 1,
      "story_type": "lead",
      "section": "safety",
      "editorial_classification": "editorial",
      "headline": "The Reasoning Was Readable. Its Important Steps Were Still Hidden",
      "slug": "the-reasoning-was-readable-its-important-steps-were-still-hidden",
      "dek": "Rollout-based advantage tests found that language-model judges could identify consequential chain-of-thought steps better than chance—but far below the experiment's noise ceiling.",
      "summary": "Rollout-based advantage tests found that language-model judges could identify consequential chain-of-thought steps better than chance—but far below the experiment's noise ceiling.",
      "body_text": "Researchers estimated each reasoning step's functional importance by measuring how much including it changed the expected final reward across Monte Carlo continuations. Capable language models beat a prevalence baseline when asked to identify high-advantage steps from the text alone, yet remained well short of the noise ceiling. Fine-tuning a step critic helped more on wrong answers than correct ones. The result does not show that reasoning text is useless; it shows that readable prose only partially reveals which step actually carries the answer. That distinction matters when traces are used for error diagnosis, process rewards or claims of interpretability.",
      "why_it_matters": "Rollout-based advantage tests found that language-model judges could identify consequential chain-of-thought steps better than chance—but far below the experiment's noise ceiling.",
      "limitations": [
        "The result does not show that reasoning text is useless; it shows that readable prose only partially reveals which step actually carries the answer."
      ],
      "importance": 10,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-001/the-reasoning-was-readable-its-important-steps-were-still-hidden",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-001.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-06-001-001",
          "text": "Rollout-based advantage tests found that language-model judges could identify consequential chain-of-thought steps better than chance—but far below the experiment's noise ceiling.",
          "source_ids": [
            "source-2026-09-06-001"
          ],
          "qualification": "The result does not show that reasoning text is useless; it shows that readable prose only partially reveals which step actually carries the answer."
        }
      ],
      "source_ids": [
        "source-2026-09-06-001"
      ],
      "tags": [
        "chain of thought",
        "interpretability",
        "process evaluation"
      ],
      "image_url": "https://themachinepress.com/issues/2026-09-06/lead-reasoning-importance.png",
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-002",
      "source_story_id": "tmp-feature-human-ai-test-time-adaptation",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 2,
      "story_type": "secondary",
      "section": "research",
      "editorial_classification": "editorial",
      "headline": "Thirty People Taught an Agent Their Own Bar",
      "slug": "thirty-people-taught-an-agent-their-own-bar",
      "dek": "Repeated feedback across writing and visual tasks became personal context, weight updates and an evolving rubric—then improved solo work after only tens of examples.",
      "summary": "Repeated feedback across writing and visual tasks became personal context, weight updates and an evolving rubric—then improved solo work after only tens of examples.",
      "body_text": "The TAHI system treats iterative human-agent work as training data for the individual rather than another sample of population preference. Across 600 writing and visual-creation tasks for 30 people, the authors report solo-task success gains of 4.5 to 20.9 percent after tens of tasks. An evolving rubric captured 16.0 to 22.3 percent more failures than rubrics produced by language models or people alone, and some improvements transferred across users by up to 8.8 percent. These results belong to the study's participants, domains and evaluation setup; they do not establish universal personalization or eliminate the need for human review.",
      "why_it_matters": "Repeated feedback across writing and visual tasks became personal context, weight updates and an evolving rubric—then improved solo work after only tens of examples.",
      "limitations": [
        "These results belong to the study's participants, domains and evaluation setup; they do not establish universal personalization or eliminate the need for human review."
      ],
      "importance": 9,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-002/thirty-people-taught-an-agent-their-own-bar",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-002.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-06-002-001",
          "text": "Repeated feedback across writing and visual tasks became personal context, weight updates and an evolving rubric—then improved solo work after only tens of examples.",
          "source_ids": [
            "source-2026-09-06-002"
          ],
          "qualification": "These results belong to the study's participants, domains and evaluation setup; they do not establish universal personalization or eliminate the need for human review."
        }
      ],
      "source_ids": [
        "source-2026-09-06-002"
      ],
      "tags": [
        "personalization",
        "test-time adaptation",
        "human feedback"
      ],
      "image_url": "https://themachinepress.com/issues/2026-09-06/feature-human-ai-adaptation.png",
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-003",
      "source_story_id": "tmp-story-temporal-self-distillation-video",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 3,
      "story_type": "dispatch",
      "section": "frontier-models",
      "editorial_classification": "editorial",
      "headline": "The Dense Video View Taught the Sparse One What Changed",
      "slug": "the-dense-video-view-taught-the-sparse-one-what-changed",
      "dek": "A model used denser temporal sampling as its own teacher, improving state tracking without labels, a separate teacher or added inference cost.",
      "summary": "A model used denser temporal sampling as its own teacher, improving state tracking without labels, a separate teacher or added inference cost.",
      "body_text": "S3T gives the same model a dense view of a clip as privileged training information and asks a sparse-view student to match its next-token distribution. On LLaVA-OneVision-2-8B, the authors report VSTAT gains from 1.74 to 2.70 points depending on configuration. Training on unlabeled synthetic clips also transferred to real video, adding 7.95 points on VSTAT-YouTube and 4.50 on MVBench Action Count. Those gains are benchmark results for the tested model, not proof of general video understanding.",
      "why_it_matters": "A model used denser temporal sampling as its own teacher, improving state tracking without labels, a separate teacher or added inference cost.",
      "limitations": [
        "Those gains are benchmark results for the tested model, not proof of general video understanding."
      ],
      "importance": 9,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-003/the-dense-video-view-taught-the-sparse-one-what-changed",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-003.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-06-003-001",
          "text": "A model used denser temporal sampling as its own teacher, improving state tracking without labels, a separate teacher or added inference cost.",
          "source_ids": [
            "source-2026-09-06-003"
          ],
          "qualification": "Those gains are benchmark results for the tested model, not proof of general video understanding."
        }
      ],
      "source_ids": [
        "source-2026-09-06-003"
      ],
      "tags": [
        "video understanding",
        "self-distillation",
        "state tracking"
      ],
      "image_url": "https://themachinepress.com/issues/2026-09-06/temporal-distillation-file.webp",
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-004",
      "source_story_id": "tmp-story-scal3r-online-reconstruction",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 4,
      "story_type": "dispatch",
      "section": "robotics",
      "editorial_classification": "editorial",
      "headline": "The Depth Stayed Sound While the Camera Pose Drifted Away",
      "slug": "the-depth-stayed-sound-while-the-camera-pose-drifted-away",
      "dek": "Scal3R froze the geometry backbone, added multi-reference pose tokens and cut average trajectory error by more than 60 percent on KITTI.",
      "summary": "Scal3R froze the geometry backbone, added multi-reference pose tokens and cut average trajectory error by more than 60 percent on KITTI.",
      "body_text": "Long online reconstructions often collapse because every pose is extrapolated from the first frame even when per-frame depth remains stable. Scal3R adds lightweight tokens—about one percent of the model's parameters—to query pose against multiple past keyframes, then closes loops with online pose-graph optimization. The authors report convergence in eight hours on one GPU and state-of-the-art results across six evaluated datasets. The evidence concerns benchmark reconstruction, not safety certification for deployed navigation.",
      "why_it_matters": "Scal3R froze the geometry backbone, added multi-reference pose tokens and cut average trajectory error by more than 60 percent on KITTI.",
      "limitations": [
        "The evidence concerns benchmark reconstruction, not safety certification for deployed navigation."
      ],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-004/the-depth-stayed-sound-while-the-camera-pose-drifted-away",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-004.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-06-004-001",
          "text": "Scal3R froze the geometry backbone, added multi-reference pose tokens and cut average trajectory error by more than 60 percent on KITTI.",
          "source_ids": [
            "source-2026-09-06-004"
          ],
          "qualification": "The evidence concerns benchmark reconstruction, not safety certification for deployed navigation."
        }
      ],
      "source_ids": [
        "source-2026-09-06-004"
      ],
      "tags": [
        "3D reconstruction",
        "camera pose",
        "loop closure"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-005",
      "source_story_id": "tmp-story-espo-prompt-optimization",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 5,
      "story_type": "dispatch",
      "section": "developer-tools",
      "editorial_classification": "editorial",
      "headline": "The Better Prompt Was Forty-Seven Percent Shorter",
      "slug": "the-better-prompt-was-forty-seven-percent-shorter",
      "dek": "ESPO clustered errors before proposing diverse fixes, then used bootstrap stability to keep prompt search from rewarding brittle gains.",
      "summary": "ESPO clustered errors before proposing diverse fixes, then used bootstrap stability to keep prompt search from rewarding brittle gains.",
      "body_text": "Across seven public NLP benchmarks, ESPO averaged 74.67 percent accuracy against 70.91 percent for GEPA while producing prompts of 1,004 rather than 1,878 characters. The method separates diagnosis, four proposal strategies and a stability-based selection stage. An ablation found that adding diversity without bootstrap selection reduced performance by 1.20 points. Cross-model gains were reported on four additional students, but the large per-task variation means the average should not be treated as a universal prompt-optimization guarantee.",
      "why_it_matters": "ESPO clustered errors before proposing diverse fixes, then used bootstrap stability to keep prompt search from rewarding brittle gains.",
      "limitations": [
        "Cross-model gains were reported on four additional students, but the large per-task variation means the average should not be treated as a universal prompt-optimization guarantee."
      ],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-005/the-better-prompt-was-forty-seven-percent-shorter",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-005.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-06-005-001",
          "text": "ESPO clustered errors before proposing diverse fixes, then used bootstrap stability to keep prompt search from rewarding brittle gains.",
          "source_ids": [
            "source-2026-09-06-005"
          ],
          "qualification": "Cross-model gains were reported on four additional students, but the large per-task variation means the average should not be treated as a universal prompt-optimization guarantee."
        }
      ],
      "source_ids": [
        "source-2026-09-06-005"
      ],
      "tags": [
        "prompt optimization",
        "evaluation",
        "stability"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-006",
      "source_story_id": "tmp-story-puffin-native-3d-world-states",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 6,
      "story_type": "dispatch",
      "section": "frontier-models",
      "editorial_classification": "editorial",
      "headline": "The World Model Kept Physics, Depth and Appearance Together",
      "slug": "the-world-model-kept-physics-depth-and-appearance-together",
      "dek": "Puffin-World represents gravity and latitude, geometry and imagery inside one multimodal generator instead of handing 3D state to offline modules.",
      "summary": "Puffin-World represents gravity and latitude, geometry and imagery inside one multimodal generator instead of handing 3D state to offline modules.",
      "body_text": "The architecture jointly models physical state, depth and appearance with a shared camera representation, then propagates dynamics into future frames. Its training collection contains 15 million vision-language-camera triplets and one million motion trajectories. The team also reports closed-loop exploration demonstrations and released code, models and datasets. The paper presents a research system and benchmark evidence; it does not establish physically reliable simulation for safety-critical decisions.",
      "why_it_matters": "Puffin-World represents gravity and latitude, geometry and imagery inside one multimodal generator instead of handing 3D state to offline modules.",
      "limitations": [
        "The paper presents a research system and benchmark evidence; it does not establish physically reliable simulation for safety-critical decisions."
      ],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-006/the-world-model-kept-physics-depth-and-appearance-together",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-006.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-06-006-001",
          "text": "Puffin-World represents gravity and latitude, geometry and imagery inside one multimodal generator instead of handing 3D state to offline modules.",
          "source_ids": [
            "source-2026-09-06-006"
          ],
          "qualification": "The paper presents a research system and benchmark evidence; it does not establish physically reliable simulation for safety-critical decisions."
        }
      ],
      "source_ids": [
        "source-2026-09-06-006"
      ],
      "tags": [
        "world models",
        "3D generation",
        "physical simulation"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-007",
      "source_story_id": "tmp-story-editvid-unified-video-editing",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 7,
      "story_type": "dispatch",
      "section": "media-creative-tools",
      "editorial_classification": "editorial",
      "headline": "One Video Editor Learned Six Kinds of Change Without Training",
      "slug": "one-video-editor-learned-six-kinds-of-change-without-training",
      "dek": "EditVid combines sparse causal memory, token correspondence and latent blending for instruction- and reference-guided edits.",
      "summary": "EditVid combines sparse causal memory, token correspondence and latent blending for instruction- and reference-guided edits.",
      "body_text": "The framework supports style transfer, attribute changes, object insertion, part edits and subject replacement without task-specific training. On FiVE, the authors report 78.16 FiVE-Acc versus 58.95 for the strongest evaluated training-free baseline, with competitive IVEBench results. A user study preferred EditVid overall in 51.8 percent of comparisons against seven methods. Those numbers reflect the chosen benchmarks and comparisons, not a blanket claim of identity-safe or artifact-free editing.",
      "why_it_matters": "EditVid combines sparse causal memory, token correspondence and latent blending for instruction- and reference-guided edits.",
      "limitations": [
        "Those numbers reflect the chosen benchmarks and comparisons, not a blanket claim of identity-safe or artifact-free editing."
      ],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-007/one-video-editor-learned-six-kinds-of-change-without-training",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-007.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-06-007-001",
          "text": "EditVid combines sparse causal memory, token correspondence and latent blending for instruction- and reference-guided edits.",
          "source_ids": [
            "source-2026-09-06-007"
          ],
          "qualification": "Those numbers reflect the chosen benchmarks and comparisons, not a blanket claim of identity-safe or artifact-free editing."
        }
      ],
      "source_ids": [
        "source-2026-09-06-007"
      ],
      "tags": [
        "video editing",
        "identity preservation",
        "generative media"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-008",
      "source_story_id": "tmp-story-small-model-declarative-ui",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 8,
      "story_type": "dispatch",
      "section": "developer-tools",
      "editorial_classification": "editorial",
      "headline": "A Four-Billion-Parameter Model Recovered Nearly All the UI Teacher's Score",
      "slug": "a-four-billion-parameter-model-recovered-nearly-all-the-ui-teacher-s-score",
      "dek": "Catalog-conditioned fine-tuning reached roughly 98 percent of teacher semantic quality and 97 percent of visual quality at far lower reported cost.",
      "summary": "Catalog-conditioned fine-tuning reached roughly 98 percent of teacher semantic quality and 97 percent of visual quality at far lower reported cost.",
      "body_text": "The study evaluates declarative interface generation, where a model selects approved components and binds data rather than writing arbitrary frontend code. Across two React and TypeScript domains, the 4B student retained nearly all measured teacher quality at more than an order of magnitude lower cost. Perturbed-catalog and constrained-ground-truth training each improved the quality-cost frontier in different ways. The result is specific to the tested component catalogs, domains, checkpoints and scoring system.",
      "why_it_matters": "Catalog-conditioned fine-tuning reached roughly 98 percent of teacher semantic quality and 97 percent of visual quality at far lower reported cost.",
      "limitations": [],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-008/a-four-billion-parameter-model-recovered-nearly-all-the-ui-teacher-s-score",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-008.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-06-008-001",
          "text": "Catalog-conditioned fine-tuning reached roughly 98 percent of teacher semantic quality and 97 percent of visual quality at far lower reported cost.",
          "source_ids": [
            "source-2026-09-06-008"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-09-06-008"
      ],
      "tags": [
        "declarative UI",
        "small models",
        "component catalogs"
      ],
      "image_url": "https://themachinepress.com/issues/2026-09-06/declarative-ui-code-file.webp",
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-009",
      "source_story_id": "tmp-story-pretraining-auxiliary-views",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 9,
      "story_type": "dispatch",
      "section": "research",
      "editorial_classification": "editorial",
      "headline": "The Model Learned a Fact Better When the Corpus Showed Another View",
      "slug": "the-model-learned-a-fact-better-when-the-corpus-showed-another-view",
      "dek": "With a fixed token budget, reformulations beat spending the same tokens on simple document repetition—even for factual recall.",
      "summary": "With a fixed token budget, reformulations beat spending the same tokens on simple document repetition—even for factual recall.",
      "body_text": "Controlled pretraining experiments found that repetition remained necessary, but reallocating some repeated tokens to auxiliary representations improved knowledge acquisition. Paraphrases helped under smaller batches, while contextual and foundational views aided learning when prior knowledge was missing. The effect did not depend on a stronger teacher generating the reformulation. The experiments isolate learning mechanisms under controlled conditions and do not prove that every synthetic rewrite improves a production corpus.",
      "why_it_matters": "With a fixed token budget, reformulations beat spending the same tokens on simple document repetition—even for factual recall.",
      "limitations": [
        "The effect did not depend on a stronger teacher generating the reformulation.",
        "The experiments isolate learning mechanisms under controlled conditions and do not prove that every synthetic rewrite improves a production corpus."
      ],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-009/the-model-learned-a-fact-better-when-the-corpus-showed-another-view",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-009.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-06-009-001",
          "text": "With a fixed token budget, reformulations beat spending the same tokens on simple document repetition—even for factual recall.",
          "source_ids": [
            "source-2026-09-06-009"
          ],
          "qualification": "The effect did not depend on a stronger teacher generating the reformulation."
        }
      ],
      "source_ids": [
        "source-2026-09-06-009"
      ],
      "tags": [
        "pretraining",
        "data diversity",
        "knowledge acquisition"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-010",
      "source_story_id": "tmp-story-probabilistic-causal-impact",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 10,
      "story_type": "dispatch",
      "section": "research",
      "editorial_classification": "editorial",
      "headline": "Actual Causality Moved Out of the Toy Model",
      "slug": "actual-causality-moved-out-of-the-toy-model",
      "dek": "Probabilistic Causal Impact turns blame and credit into a Monte Carlo estimation problem over an explicit causal model.",
      "summary": "Probabilistic Causal Impact turns blame and credit into a Monte Carlo estimation problem over an explicit causal model.",
      "body_text": "PCI combines ideas from actual causality and probabilities of necessity and sufficiency while letting users define candidate explanations, counterfactual values and scores. The authors test consistency against exact causal verdicts, scale the method in synthetic systems and apply it to a deployed causal model trained on millions of points. The framework produces graded, causally structured explanations rather than feature attribution alone. Results still depend on the assumed causal graph and counterfactual distributions; computation cannot repair a misspecified model.",
      "why_it_matters": "Probabilistic Causal Impact turns blame and credit into a Monte Carlo estimation problem over an explicit causal model.",
      "limitations": [
        "Results still depend on the assumed causal graph and counterfactual distributions; computation cannot repair a misspecified model."
      ],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-010/actual-causality-moved-out-of-the-toy-model",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-010.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-06-010-001",
          "text": "Probabilistic Causal Impact turns blame and credit into a Monte Carlo estimation problem over an explicit causal model.",
          "source_ids": [
            "source-2026-09-06-010"
          ],
          "qualification": "Results still depend on the assumed causal graph and counterfactual distributions; computation cannot repair a misspecified model."
        }
      ],
      "source_ids": [
        "source-2026-09-06-010"
      ],
      "tags": [
        "causal inference",
        "explanation",
        "Monte Carlo"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-011",
      "source_story_id": "tmp-story-last-translation-benchmark",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 11,
      "story_type": "dispatch",
      "section": "benchmarks-evals",
      "editorial_classification": "editorial",
      "headline": "The Translation Benchmark Collects Only Examples That Still Break Models",
      "slug": "the-translation-benchmark-collects-only-examples-that-still-break-models",
      "dek": "A live, peer-reviewed dataset pairs difficult multimodal translations with handcrafted rules that say exactly what failure looks like.",
      "summary": "A live, peer-reviewed dataset pairs difficult multimodal translations with handcrafted rules that say exactly what failure looks like.",
      "body_text": "The Last Translation Benchmark starts from human-authored examples that defeat leading systems rather than a static set approaching saturation. Each text, image, audio or video case includes verification rules for concrete errors, aiming to make evaluation more reproducible and actionable than a single automatic score. Version one includes accepted contributions before September 1 and is designed to keep growing. Its value will depend on contribution quality, coverage and sustained review; it is not itself evidence that translation progress has stopped.",
      "why_it_matters": "A live, peer-reviewed dataset pairs difficult multimodal translations with handcrafted rules that say exactly what failure looks like.",
      "limitations": [
        "Its value will depend on contribution quality, coverage and sustained review; it is not itself evidence that translation progress has stopped."
      ],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-011/the-translation-benchmark-collects-only-examples-that-still-break-models",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-011.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-06-011-001",
          "text": "A live, peer-reviewed dataset pairs difficult multimodal translations with handcrafted rules that say exactly what failure looks like.",
          "source_ids": [
            "source-2026-09-06-011"
          ],
          "qualification": "Its value will depend on contribution quality, coverage and sustained review; it is not itself evidence that translation progress has stopped."
        }
      ],
      "source_ids": [
        "source-2026-09-06-011"
      ],
      "tags": [
        "machine translation",
        "benchmark",
        "verification"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-012",
      "source_story_id": "tmp-story-one-example-on-policy-distillation",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 12,
      "story_type": "dispatch",
      "section": "frontier-models",
      "editorial_classification": "editorial",
      "headline": "One Training Query Reached Seventy-One Percent of the Teacher's States",
      "slug": "one-training-query-reached-seventy-one-percent-of-the-teacher-s-states",
      "dek": "On-policy distillation kept improving for hundreds of steps even when the student repeatedly learned from a single prompt.",
      "summary": "On-policy distillation kept improving for hundreds of steps even when the student repeatedly learned from a single prompt.",
      "body_text": "A single query's rollouts visited 71.5 percent of the states reached by full-data training, mostly within the first 100 steps. Sixteen semantically diverse queries reached 98.9 percent coverage and matched full-data gains, while content-light and off-domain prompts approached the real-query baseline. The authors argue that on-policy distillation is data-overfed but algorithm-starved: rollouts expose broad supervision quickly, then alignment absorbs it slowly. The finding is bounded to the tested tasks, teachers and model families.",
      "why_it_matters": "On-policy distillation kept improving for hundreds of steps even when the student repeatedly learned from a single prompt.",
      "limitations": [],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-012/one-training-query-reached-seventy-one-percent-of-the-teacher-s-states",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-012.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-06-012-001",
          "text": "On-policy distillation kept improving for hundreds of steps even when the student repeatedly learned from a single prompt.",
          "source_ids": [
            "source-2026-09-06-012"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-09-06-012"
      ],
      "tags": [
        "distillation",
        "state coverage",
        "post-training"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-013",
      "source_story_id": "tmp-story-para-pipe-soc-parallelism",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 13,
      "story_type": "dispatch",
      "section": "chips-infrastructure",
      "editorial_classification": "editorial",
      "headline": "The Edge Chip Stopped Choosing Between a Pipeline and Parallel Work",
      "slug": "the-edge-chip-stopped-choosing-between-a-pipeline-and-parallel-work",
      "dek": "Para-Pipe maps operator concurrency within and across stages, producing Pareto choices for latency, throughput and energy on heterogeneous SoCs.",
      "summary": "Para-Pipe maps operator concurrency within and across stages, producing Pareto choices for latency, throughput and energy on heterogeneous SoCs.",
      "body_text": "The framework searches how a neural graph should share work across big and little CPU cores, a GPU, DSPs and a dedicated accelerator. On one Amlogic system, throughput-optimized configurations improved average energy efficiency by 11.0 percent over pure pipelining and 23.3 percent over non-pipelined parallel execution. A second automotive-class platform supplied another heterogeneous test. These are measurements on the authors' graphs and devices, not general efficiency guarantees for all edge workloads.",
      "why_it_matters": "Para-Pipe maps operator concurrency within and across stages, producing Pareto choices for latency, throughput and energy on heterogeneous SoCs.",
      "limitations": [
        "These are measurements on the authors' graphs and devices, not general efficiency guarantees for all edge workloads."
      ],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-013/the-edge-chip-stopped-choosing-between-a-pipeline-and-parallel-work",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-013.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-06-013-001",
          "text": "Para-Pipe maps operator concurrency within and across stages, producing Pareto choices for latency, throughput and energy on heterogeneous SoCs.",
          "source_ids": [
            "source-2026-09-06-013"
          ],
          "qualification": "These are measurements on the authors' graphs and devices, not general efficiency guarantees for all edge workloads."
        }
      ],
      "source_ids": [
        "source-2026-09-06-013"
      ],
      "tags": [
        "edge AI",
        "operator parallelism",
        "energy efficiency"
      ],
      "image_url": "https://themachinepress.com/issues/2026-09-06/soc-hardware-file.webp",
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-014",
      "source_story_id": "tmp-story-cnot-synthesis-np-hard",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 14,
      "story_type": "dispatch",
      "section": "research",
      "editorial_classification": "editorial",
      "headline": "Even the Plain CNOT Circuit Problem Is NP-Hard",
      "slug": "even-the-plain-cnot-circuit-problem-is-np-hard",
      "dek": "A reduction from Hamiltonian paths closes the complexity question for exact synthesis with labeled qubits, all-to-all links and no ancillas.",
      "summary": "A reduction from Hamiltonian paths closes the complexity question for exact synthesis with labeled qubits, all-to-all links and no ancillas.",
      "body_text": "Earlier hardness proofs needed restricted connectivity, encoded inputs or extra intermediate variables. The new proof uses recorder qubits to force required intermediate visits into the final parity transformation, reducing a grid-graph Hamiltonian path to the vanilla synthesis problem. The decision form is NP-complete and optimization is NP-hard, with consequences for related shortest-word, Cayley-graph distance and XOR-program problems. Complexity hardness describes worst-case computation; it does not say useful circuits cannot be optimized in practice.",
      "why_it_matters": "A reduction from Hamiltonian paths closes the complexity question for exact synthesis with labeled qubits, all-to-all links and no ancillas.",
      "limitations": [
        "Complexity hardness describes worst-case computation; it does not say useful circuits cannot be optimized in practice."
      ],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-014/even-the-plain-cnot-circuit-problem-is-np-hard",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-014.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-06-014-001",
          "text": "A reduction from Hamiltonian paths closes the complexity question for exact synthesis with labeled qubits, all-to-all links and no ancillas.",
          "source_ids": [
            "source-2026-09-06-014"
          ],
          "qualification": "Complexity hardness describes worst-case computation; it does not say useful circuits cannot be optimized in practice."
        }
      ],
      "source_ids": [
        "source-2026-09-06-014"
      ],
      "tags": [
        "quantum circuits",
        "complexity",
        "CNOT synthesis"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-026",
      "source_story_id": "tmp-story-generative-image-identity-benchmark",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 15,
      "story_type": "dispatch",
      "section": "benchmarks-evals",
      "editorial_classification": "editorial",
      "headline": "Image Quality Rose While the Subject's Identity Drifted",
      "slug": "image-quality-rose-while-the-subject-s-identity-drifted",
      "dek": "A benchmark separates fidelity from polish across generation, editing, restoration and multi-subject scenes, then tests identity as persistent knowledge.",
      "summary": "A benchmark separates fidelity from polish across generation, editing, restoration and multi-subject scenes, then tests identity as persistent knowledge.",
      "body_text": "The study compares identity supplied in prompt context, encoded in subject-specific parameters and maintained through a persistent identity layer. Drift worsened during repeated edits, at small subject scales, under severe restoration and when several subjects shared a scene. The persistent representation improved identity scores across tested foundation models while preserving comparable instruction adherence and perceptual quality. The benchmark is produced alongside one of the compared approaches, so the result should be read as reported evaluation rather than neutral product certification.",
      "why_it_matters": "A benchmark separates fidelity from polish across generation, editing, restoration and multi-subject scenes, then tests identity as persistent knowledge.",
      "limitations": [],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-026/image-quality-rose-while-the-subject-s-identity-drifted",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-026.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-06-026-001",
          "text": "A benchmark separates fidelity from polish across generation, editing, restoration and multi-subject scenes, then tests identity as persistent knowledge.",
          "source_ids": [
            "source-2026-09-06-015"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-09-06-015"
      ],
      "tags": [
        "image generation",
        "identity preservation",
        "benchmark"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-027",
      "source_story_id": "tmp-story-nlip-agent-standard",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 16,
      "story_type": "dispatch",
      "section": "open-source",
      "editorial_classification": "editorial",
      "headline": "Agent Interoperability Got a Natural-Language Envelope",
      "slug": "agent-interoperability-got-a-natural-language-envelope",
      "dek": "The Ecma-standardized NLIP defines a semantic message layer that can travel over HTTP, WebSocket or AMQP while adapting to other agent protocols.",
      "summary": "The Ecma-standardized NLIP defines a semantic message layer that can travel over HTTP, WebSocket or AMQP while adapting to other agent protocols.",
      "body_text": "NLIP wraps interaction in a common application-layer message model rather than replacing transports or every tool protocol. The paper describes bindings, security considerations, a reference implementation, representative applications and early adoption signals, plus its relationship to MCP and A2A. Standardization gives implementers a shared specification; it does not prove broad deployment, automatic semantic agreement or secure behavior by every conforming agent.",
      "why_it_matters": "The Ecma-standardized NLIP defines a semantic message layer that can travel over HTTP, WebSocket or AMQP while adapting to other agent protocols.",
      "limitations": [
        "Standardization gives implementers a shared specification; it does not prove broad deployment, automatic semantic agreement or secure behavior by every conforming agent."
      ],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-027/agent-interoperability-got-a-natural-language-envelope",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-027.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-06-027-001",
          "text": "The Ecma-standardized NLIP defines a semantic message layer that can travel over HTTP, WebSocket or AMQP while adapting to other agent protocols.",
          "source_ids": [
            "source-2026-09-06-016"
          ],
          "qualification": "Standardization gives implementers a shared specification; it does not prove broad deployment, automatic semantic agreement or secure behavior by every conforming agent."
        }
      ],
      "source_ids": [
        "source-2026-09-06-016"
      ],
      "tags": [
        "agent protocols",
        "interoperability",
        "standards"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-015",
      "source_story_id": "tmp-sidebar-latentstream-video-memory",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 17,
      "story_type": "ticker",
      "section": "frontier-models",
      "editorial_classification": "editorial",
      "headline": "Streaming Video Memory Moved Inside the Model",
      "slug": "streaming-video-memory-moved-inside-the-model",
      "dek": "LatentStream consolidates short-, mid- and long-term evidence, then internalizes retrieved history into a fixed-length latent state.",
      "summary": "LatentStream consolidates short-, mid- and long-term evidence, then internalizes retrieved history into a fixed-length latent state.",
      "body_text": "The framework uses adaptive hierarchical consolidation and confidence-guided optimization under a bounded memory budget. The authors report state-of-the-art results on tested online and offline video benchmarks.",
      "why_it_matters": "LatentStream consolidates short-, mid- and long-term evidence, then internalizes retrieved history into a fixed-length latent state.",
      "limitations": [],
      "importance": 7,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-015/streaming-video-memory-moved-inside-the-model",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-015.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-06-015-001",
          "text": "LatentStream consolidates short-, mid- and long-term evidence, then internalizes retrieved history into a fixed-length latent state.",
          "source_ids": [
            "source-2026-09-06-017"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-09-06-017"
      ],
      "tags": [
        "streaming video",
        "latent memory"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-016",
      "source_story_id": "tmp-sidebar-atiba-paper-integrity",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 18,
      "story_type": "ticker",
      "section": "developer-tools",
      "editorial_classification": "editorial",
      "headline": "The Paper Checker Refused Evidence It Could Not Find Verbatim",
      "slug": "the-paper-checker-refused-evidence-it-could-not-find-verbatim",
      "dek": "ATIBA grounds venue rules, citations and reporting checks in source pages and manuscript text, then discards missing evidence quotes.",
      "summary": "ATIBA grounds venue rules, citations and reporting checks in source pages and manuscript text, then discards missing evidence quotes.",
      "body_text": "A 13-person moderated study found 69 to 92 percent agreement across six usefulness items, averaging 85 percent. Objective accuracy remains unmeasured, which the authors state explicitly.",
      "why_it_matters": "ATIBA grounds venue rules, citations and reporting checks in source pages and manuscript text, then discards missing evidence quotes.",
      "limitations": [],
      "importance": 7,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-016/the-paper-checker-refused-evidence-it-could-not-find-verbatim",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-016.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-06-016-001",
          "text": "ATIBA grounds venue rules, citations and reporting checks in source pages and manuscript text, then discards missing evidence quotes.",
          "source_ids": [
            "source-2026-09-06-018"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-09-06-018"
      ],
      "tags": [
        "research integrity",
        "grounding"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-017",
      "source_story_id": "tmp-sidebar-deception-causal-mechanisms",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 19,
      "story_type": "ticker",
      "section": "safety",
      "editorial_classification": "editorial",
      "headline": "Deceptive-Looking Output Did Not Always Mean a Deceptive Mechanism",
      "slug": "deceptive-looking-output-did-not-always-mean-a-deceptive-mechanism",
      "dek": "Controlled guessing and trading tasks separate misleading behavior, recipient-state sensitivity and claims of agency.",
      "summary": "Controlled guessing and trading tasks separate misleading behavior, recipient-state sensitivity and claims of agency.",
      "body_text": "Interventions found both false positives—deceptive-looking behavior without the proposed mechanism—and cases where recipient information causally changed deceptive preference. Even the latter does not establish model agency.",
      "why_it_matters": "Controlled guessing and trading tasks separate misleading behavior, recipient-state sensitivity and claims of agency.",
      "limitations": [
        "Even the latter does not establish model agency."
      ],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-017/deceptive-looking-output-did-not-always-mean-a-deceptive-mechanism",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-017.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-06-017-001",
          "text": "Controlled guessing and trading tasks separate misleading behavior, recipient-state sensitivity and claims of agency.",
          "source_ids": [
            "source-2026-09-06-019"
          ],
          "qualification": "Even the latter does not establish model agency."
        }
      ],
      "source_ids": [
        "source-2026-09-06-019"
      ],
      "tags": [
        "deception",
        "causal analysis"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-018",
      "source_story_id": "tmp-sidebar-tokenmatch-mesh-correspondence",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 20,
      "story_type": "ticker",
      "section": "research",
      "editorial_classification": "editorial",
      "headline": "Curvature Turned a Mesh Into Tokens",
      "slug": "curvature-turned-a-mesh-into-tokens",
      "dek": "TokenMatch learned partial and full 3D correspondences with adaptive patches and sub-second feed-forward inference.",
      "summary": "TokenMatch learned partial and full 3D correspondences with adaptive patches and sub-second feed-forward inference.",
      "body_text": "Trained only on a partial-shape dataset, the transformer generalized to full-shape benchmarks without fine-tuning and reported strong geodesic-error and overlap results across six suites.",
      "why_it_matters": "TokenMatch learned partial and full 3D correspondences with adaptive patches and sub-second feed-forward inference.",
      "limitations": [
        "Trained only on a partial-shape dataset, the transformer generalized to full-shape benchmarks without fine-tuning and reported strong geodesic-error and overlap results across six suites."
      ],
      "importance": 7,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-018/curvature-turned-a-mesh-into-tokens",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-018.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-06-018-001",
          "text": "TokenMatch learned partial and full 3D correspondences with adaptive patches and sub-second feed-forward inference.",
          "source_ids": [
            "source-2026-09-06-020"
          ],
          "qualification": "Trained only on a partial-shape dataset, the transformer generalized to full-shape benchmarks without fine-tuning and reported strong geodesic-error and overlap results across six suites."
        }
      ],
      "source_ids": [
        "source-2026-09-06-020"
      ],
      "tags": [
        "3D geometry",
        "mesh correspondence"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-019",
      "source_story_id": "tmp-sidebar-ntep-tool-evidence-paths",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 21,
      "story_type": "ticker",
      "section": "benchmarks-evals",
      "editorial_classification": "editorial",
      "headline": "Every Tool Call Had to Advance an Evidence Goal",
      "slug": "every-tool-call-had-to-advance-an-evidence-goal",
      "dek": "NTEP rewards an agentic vision model for seeking necessary external evidence and penalizes repeated searches for an already satisfied need.",
      "summary": "NTEP rewards an agentic vision model for seeking necessary external evidence and penalizes repeated searches for an already satisfied need.",
      "body_text": "The 8B implementation improved search accuracy and tool efficiency across seven image-grounded benchmarks. The result supports finer-grained supervision for tested crop, image-search and text-search tools.",
      "why_it_matters": "NTEP rewards an agentic vision model for seeking necessary external evidence and penalizes repeated searches for an already satisfied need.",
      "limitations": [],
      "importance": 7,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-019/every-tool-call-had-to-advance-an-evidence-goal",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-019.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-06-019-001",
          "text": "NTEP rewards an agentic vision model for seeking necessary external evidence and penalizes repeated searches for an already satisfied need.",
          "source_ids": [
            "source-2026-09-06-021"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-09-06-021"
      ],
      "tags": [
        "tool use",
        "vision-language models"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-020",
      "source_story_id": "invention-hangprinter",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 22,
      "story_type": "invention_desk",
      "section": "invention-desk",
      "editorial_classification": "editorial",
      "headline": "Hangprinter",
      "slug": "hangprinter",
      "dek": "Suspends a print head from tensioned lines anchored around a room, replacing a rigid gantry with cable geometry so an open RepRap can work across an unusually large build space.",
      "summary": "Suspends a print head from tensioned lines anchored around a room, replacing a rigid gantry with cable geometry so an open RepRap can work across an unusually large build space.",
      "body_text": "Suspends a print head from tensioned lines anchored around a room, replacing a rigid gantry with cable geometry so an open RepRap can work across an unusually large build space.",
      "why_it_matters": "An independent builder is turning an improbable idea into a working project.",
      "limitations": [
        "A Desk Pick is an editorial selection, not a product endorsement."
      ],
      "importance": null,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-020/hangprinter",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-020.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [],
      "source_ids": [],
      "tags": [
        "independent-builders",
        "desk-pick",
        "prototype"
      ],
      "image_url": "https://themachinepress.com/issues/2026-09-06/invention-desk/hangprinter.png",
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-021",
      "source_story_id": "invention-precious-plastic",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 23,
      "story_type": "invention_desk",
      "section": "invention-desk",
      "editorial_classification": "editorial",
      "headline": "Precious Plastic",
      "slug": "precious-plastic",
      "dek": "Publishes replicable shredders, presses, workspace plans, and shared know-how so small local teams can sort waste plastic and turn it into reusable flakes and sheet material.",
      "summary": "Publishes replicable shredders, presses, workspace plans, and shared know-how so small local teams can sort waste plastic and turn it into reusable flakes and sheet material.",
      "body_text": "Publishes replicable shredders, presses, workspace plans, and shared know-how so small local teams can sort waste plastic and turn it into reusable flakes and sheet material.",
      "why_it_matters": "An independent builder is turning an improbable idea into a working project.",
      "limitations": [
        "A Desk Pick is an editorial selection, not a product endorsement."
      ],
      "importance": null,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-021/precious-plastic",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-021.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [],
      "source_ids": [],
      "tags": [
        "independent-builders",
        "desk-pick",
        "released"
      ],
      "image_url": "https://themachinepress.com/issues/2026-09-06/invention-desk/precious-plastic.png",
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-022",
      "source_story_id": "invention-watchy",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 24,
      "story_type": "invention_desk",
      "section": "invention-desk",
      "editorial_classification": "editorial",
      "headline": "Watchy",
      "slug": "watchy",
      "dek": "Pairs a square e-paper display with an ESP32-S3 and publishes the hardware, software, documentation, and case files so owners can build and program their own watch faces.",
      "summary": "Pairs a square e-paper display with an ESP32-S3 and publishes the hardware, software, documentation, and case files so owners can build and program their own watch faces.",
      "body_text": "Pairs a square e-paper display with an ESP32-S3 and publishes the hardware, software, documentation, and case files so owners can build and program their own watch faces.",
      "why_it_matters": "An independent builder is turning an improbable idea into a working project.",
      "limitations": [
        "A Desk Pick is an editorial selection, not a product endorsement."
      ],
      "importance": null,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-022/watchy",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-022.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [],
      "source_ids": [],
      "tags": [
        "independent-builders",
        "desk-pick",
        "released"
      ],
      "image_url": "https://themachinepress.com/issues/2026-09-06/invention-desk/watchy.png",
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-023",
      "source_story_id": "invention-ploopy-classic-2",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 25,
      "story_type": "invention_desk",
      "section": "invention-desk",
      "editorial_classification": "editorial",
      "headline": "Ploopy Classic 2",
      "slug": "ploopy-classic-2",
      "dek": "Turns a desktop trackball into an inspectable kit by publishing its mechanical and electrical design files, assembly documentation, and programmable QMK firmware.",
      "summary": "Turns a desktop trackball into an inspectable kit by publishing its mechanical and electrical design files, assembly documentation, and programmable QMK firmware.",
      "body_text": "Turns a desktop trackball into an inspectable kit by publishing its mechanical and electrical design files, assembly documentation, and programmable QMK firmware.",
      "why_it_matters": "An independent builder is turning an improbable idea into a working project.",
      "limitations": [
        "A Desk Pick is an editorial selection, not a product endorsement."
      ],
      "importance": null,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-023/ploopy-classic-2",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-023.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [],
      "source_ids": [],
      "tags": [
        "independent-builders",
        "desk-pick",
        "released"
      ],
      "image_url": "https://themachinepress.com/issues/2026-09-06/invention-desk/ploopy-classic-2.png",
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-024",
      "source_story_id": "invention-sponsored-house-example",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 26,
      "story_type": "invention_desk",
      "section": "invention-desk",
      "editorial_classification": "house_example",
      "headline": "The First Paid Slot",
      "slug": "the-first-paid-slot",
      "dek": "A transparent preview of paid placement with one verified link and no claim of endorsement.",
      "summary": "A transparent preview of paid placement with one verified link and no claim of endorsement.",
      "body_text": "A transparent preview of paid placement with one verified link and no claim of endorsement.\n\nHouse example - no advertiser paid. Payment will buy placement, never endorsement.",
      "why_it_matters": "This placement explains how builders can appear in The Invention Desk without purchasing editorial endorsement.",
      "limitations": [
        "House example - no advertiser paid. Payment will buy placement, never endorsement."
      ],
      "importance": null,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-024/the-first-paid-slot",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-024.json",
      "first_published_at": "2026-07-10T09:00:00.000-04:00",
      "modified_at": "2026-07-10T09:00:00.000-04:00",
      "content_status": "carried_over",
      "is_carryover": true,
      "carryover_reason": "invention_desk_seven_day_placement",
      "key_claims": [],
      "source_ids": [],
      "tags": [
        "independent-builders",
        "sponsored-project",
        "open"
      ],
      "image_url": "https://themachinepress.com/issues/2026-09-06/invention-desk/sponsored-project.webp",
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-06-025",
      "source_story_id": "invention-placement-cta",
      "edition_id": "mp-2026-09-06-morning-0059",
      "edition_url": "https://themachinepress.com/edition/2026-09-06",
      "position": 27,
      "story_type": "invention_desk",
      "section": "invention-desk",
      "editorial_classification": "house_example",
      "headline": "Put Your Project on the Desk",
      "slug": "put-your-project-on-the-desk",
      "dek": "One manually reviewed placement stays active for seven days and remains separate from Desk Picks.",
      "summary": "One manually reviewed placement stays active for seven days and remains separate from Desk Picks.",
      "body_text": "One manually reviewed placement stays active for seven days and remains separate from Desk Picks.\n\nManual intake only. Payment buys placement, never endorsement, and every submission is reviewed.",
      "why_it_matters": "This placement explains how builders can appear in The Invention Desk without purchasing editorial endorsement.",
      "limitations": [
        "Manual intake only. Payment buys placement, never endorsement, and every submission is reviewed."
      ],
      "importance": null,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-06-025/put-your-project-on-the-desk",
      "json_url": "https://themachinepress.com/story/mp-2026-09-06-025.json",
      "first_published_at": "2026-07-10T09:00:00.000-04:00",
      "modified_at": "2026-07-10T09:00:00.000-04:00",
      "content_status": "carried_over",
      "is_carryover": true,
      "carryover_reason": "invention_desk_seven_day_placement",
      "key_claims": [],
      "source_ids": [],
      "tags": [
        "independent-builders",
        "cta",
        "open"
      ],
      "image_url": "https://themachinepress.com/issues/2026-09-06/invention-desk/put-your-project-on-the-desk.webp",
      "corrections": []
    }
  ],
  "sources": [
    {
      "source_id": "source-2026-09-06-001",
      "title": "arXiv preprint 2609.04194",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.04194",
      "canonical_url": "https://arxiv.org/abs/2609.04194",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-03T13:59:08.000-04:00",
      "accessed_at": "2026-09-06T08:20:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-06-001-001"
      ]
    },
    {
      "source_id": "source-2026-09-06-002",
      "title": "arXiv preprint 2609.04141",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.04141",
      "canonical_url": "https://arxiv.org/abs/2609.04141",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-03T13:33:18.000-04:00",
      "accessed_at": "2026-09-06T08:20:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-06-002-001"
      ]
    },
    {
      "source_id": "source-2026-09-06-003",
      "title": "arXiv preprint 2609.04203",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.04203",
      "canonical_url": "https://arxiv.org/abs/2609.04203",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-03T13:59:55.000-04:00",
      "accessed_at": "2026-09-06T08:20:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-06-003-001"
      ]
    },
    {
      "source_id": "source-2026-09-06-004",
      "title": "arXiv preprint 2609.04201",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.04201",
      "canonical_url": "https://arxiv.org/abs/2609.04201",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-03T13:59:53.000-04:00",
      "accessed_at": "2026-09-06T08:20:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-06-004-001"
      ]
    },
    {
      "source_id": "source-2026-09-06-005",
      "title": "arXiv preprint 2609.04197",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.04197",
      "canonical_url": "https://arxiv.org/abs/2609.04197",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-03T13:59:37.000-04:00",
      "accessed_at": "2026-09-06T08:20:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-06-005-001"
      ]
    },
    {
      "source_id": "source-2026-09-06-006",
      "title": "arXiv preprint 2609.04196",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.04196",
      "canonical_url": "https://arxiv.org/abs/2609.04196",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-03T13:59:13.000-04:00",
      "accessed_at": "2026-09-06T08:20:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-06-006-001"
      ]
    },
    {
      "source_id": "source-2026-09-06-007",
      "title": "arXiv preprint 2609.04190",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.04190",
      "canonical_url": "https://arxiv.org/abs/2609.04190",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-03T13:59:01.000-04:00",
      "accessed_at": "2026-09-06T08:20:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-06-007-001"
      ]
    },
    {
      "source_id": "source-2026-09-06-008",
      "title": "arXiv preprint 2609.04184",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.04184",
      "canonical_url": "https://arxiv.org/abs/2609.04184",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-03T13:58:09.000-04:00",
      "accessed_at": "2026-09-06T08:20:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-06-008-001"
      ]
    },
    {
      "source_id": "source-2026-09-06-009",
      "title": "arXiv preprint 2609.04180",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.04180",
      "canonical_url": "https://arxiv.org/abs/2609.04180",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-03T13:57:02.000-04:00",
      "accessed_at": "2026-09-06T08:20:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-06-009-001"
      ]
    },
    {
      "source_id": "source-2026-09-06-010",
      "title": "arXiv preprint 2609.04177",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.04177",
      "canonical_url": "https://arxiv.org/abs/2609.04177",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-03T13:55:43.000-04:00",
      "accessed_at": "2026-09-06T08:20:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-06-010-001"
      ]
    },
    {
      "source_id": "source-2026-09-06-011",
      "title": "arXiv preprint 2609.04173",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.04173",
      "canonical_url": "https://arxiv.org/abs/2609.04173",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-03T13:54:45.000-04:00",
      "accessed_at": "2026-09-06T08:20:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-06-011-001"
      ]
    },
    {
      "source_id": "source-2026-09-06-012",
      "title": "arXiv preprint 2609.04172",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.04172",
      "canonical_url": "https://arxiv.org/abs/2609.04172",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-03T13:54:38.000-04:00",
      "accessed_at": "2026-09-06T08:20:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-06-012-001"
      ]
    },
    {
      "source_id": "source-2026-09-06-013",
      "title": "arXiv preprint 2609.04168",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.04168",
      "canonical_url": "https://arxiv.org/abs/2609.04168",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-03T13:53:44.000-04:00",
      "accessed_at": "2026-09-06T08:20:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-06-013-001"
      ]
    },
    {
      "source_id": "source-2026-09-06-014",
      "title": "arXiv preprint 2609.04160",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.04160",
      "canonical_url": "https://arxiv.org/abs/2609.04160",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-03T13:49:16.000-04:00",
      "accessed_at": "2026-09-06T08:20:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-06-014-001"
      ]
    },
    {
      "source_id": "source-2026-09-06-015",
      "title": "arXiv preprint 2609.04151",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.04151",
      "canonical_url": "https://arxiv.org/abs/2609.04151",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-03T13:44:25.000-04:00",
      "accessed_at": "2026-09-06T08:20:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-06-026-001"
      ]
    },
    {
      "source_id": "source-2026-09-06-016",
      "title": "arXiv preprint 2609.04135",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.04135",
      "canonical_url": "https://arxiv.org/abs/2609.04135",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-03T13:30:17.000-04:00",
      "accessed_at": "2026-09-06T08:20:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-06-027-001"
      ]
    },
    {
      "source_id": "source-2026-09-06-017",
      "title": "arXiv preprint 2609.04131",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.04131",
      "canonical_url": "https://arxiv.org/abs/2609.04131",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-03T13:28:14.000-04:00",
      "accessed_at": "2026-09-06T08:20:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-06-015-001"
      ]
    },
    {
      "source_id": "source-2026-09-06-018",
      "title": "arXiv preprint 2609.04123",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.04123",
      "canonical_url": "https://arxiv.org/abs/2609.04123",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-03T13:24:15.000-04:00",
      "accessed_at": "2026-09-06T08:20:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-06-016-001"
      ]
    },
    {
      "source_id": "source-2026-09-06-019",
      "title": "arXiv preprint 2609.04166",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.04166",
      "canonical_url": "https://arxiv.org/abs/2609.04166",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-03T13:52:19.000-04:00",
      "accessed_at": "2026-09-06T08:20:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-06-017-001"
      ]
    },
    {
      "source_id": "source-2026-09-06-020",
      "title": "arXiv preprint 2609.04202",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.04202",
      "canonical_url": "https://arxiv.org/abs/2609.04202",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-03T13:59:55.000-04:00",
      "accessed_at": "2026-09-06T08:20:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-06-018-001"
      ]
    },
    {
      "source_id": "source-2026-09-06-021",
      "title": "arXiv preprint 2609.03493",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.03493",
      "canonical_url": "https://arxiv.org/abs/2609.03493",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-03T03:52:56.000-04:00",
      "accessed_at": "2026-09-06T08:20:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-06-019-001"
      ]
    }
  ],
  "corrections": [],
  "publisher": {
    "name": "The Machine Press",
    "url": "https://themachinepress.com",
    "description": "A daily newspaper for the age of artificial intelligence."
  }
}
