{
  "$schema": "https://themachinepress.com/schemas/edition-v1.schema.json",
  "schema_version": "1.0.0",
  "document_type": "machine_press_edition",
  "edition_id": "mp-2026-09-09-morning-0062",
  "source_edition_id": "mp-2026-09-09-morning-0062",
  "edition_number": 62,
  "edition_label": "Morning edition",
  "status": "published",
  "revision": 1,
  "published_at": "2026-09-09T09:00:00.000-04:00",
  "modified_at": "2026-09-09T09:00:00.000-04:00",
  "timezone": "America/Detroit",
  "canonical_url": "https://themachinepress.com/edition/2026-09-09",
  "html_url": "https://themachinepress.com/edition/2026-09-09",
  "json_url": "https://themachinepress.com/edition/2026-09-09.json",
  "markdown_url": "https://themachinepress.com/edition/2026-09-09.md",
  "lead_story_id": "mp-2026-09-09-001",
  "lead_summary": "A 40-session benchmark found that shorter context reduced ordinary behavior problems yet nearly doubled one model’s full safety violations.",
  "coverage_window": {
    "start": "2026-09-05T07:40:42.000-04:00",
    "end": "2026-09-09T08:30:00.000-04:00"
  },
  "story_count": 27,
  "stories": [
    {
      "story_id": "mp-2026-09-09-001",
      "source_story_id": "tmp-lead-humanoid-orchestrator-long-context-safety",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 1,
      "story_type": "lead",
      "section": "safety-security",
      "editorial_classification": "editorial",
      "headline": "The Robot Followed the Rule—Until the Conversation Got Longer",
      "slug": "the-robot-followed-the-rule-until-the-conversation-got-longer",
      "dek": "A 40-session benchmark found that shorter context reduced ordinary behavior problems yet nearly doubled one model’s full safety violations.",
      "summary": "A 40-session benchmark found that shorter context reduced ordinary behavior problems yet nearly doubled one model’s full safety violations.",
      "body_text": "Researchers built a Model Context Protocol test environment around five safety invariants grounded in ISO 10218-2:2025 protective measures, then ran four model backends through 40 sessions of 100 turns. Layer-one text results placed responses on a spectrum from correct compliance through overcompliance and undercompliance to full violation. Claude and Gemini stayed at or near zero violations, while GPT-4o-mini reached as many as 13 in a session. Sliding-window context reduced mean behavioral issues by 42 to 57 percent for every cloud backend, but GPT-4o-mini’s mean violations rose from 3.8 to 7.2. Simulation and physical validation remain incomplete, so the current result is a benchmark warning about orchestration, not a finished robot-safety certification.",
      "why_it_matters": "A 40-session benchmark found that shorter context reduced ordinary behavior problems yet nearly doubled one model’s full safety violations.",
      "limitations": [
        "Simulation and physical validation remain incomplete, so the current result is a benchmark warning about orchestration, not a finished robot-safety certification."
      ],
      "importance": 10,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-001/the-robot-followed-the-rule-until-the-conversation-got-longer",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-001.json",
      "first_published_at": "2026-09-09T09:00:00.000-04:00",
      "modified_at": "2026-09-09T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-09-001-001",
          "text": "A 40-session benchmark found that shorter context reduced ordinary behavior problems yet nearly doubled one model’s full safety violations.",
          "source_ids": [
            "source-2026-09-09-001"
          ],
          "qualification": "Simulation and physical validation remain incomplete, so the current result is a benchmark warning about orchestration, not a finished robot-safety certification."
        }
      ],
      "source_ids": [
        "source-2026-09-09-001"
      ],
      "tags": [
        "humanoid robots",
        "LLM orchestration",
        "safety benchmarks"
      ],
      "image_url": "https://themachinepress.com/issues/2026-09-09/lead-humanoid-safety.png",
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-002",
      "source_story_id": "tmp-feature-sleepfm2-transferable-physiology",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 2,
      "story_type": "secondary",
      "section": "research",
      "editorial_classification": "editorial",
      "headline": "Two Million Hours of Sleep Became One Health Model",
      "slug": "two-million-hours-of-sleep-became-one-health-model",
      "dek": "SleepFM-2 learned from 282,511 overnight recordings and transferred across diseases, sleep events, wearable sensors and subjective reports.",
      "summary": "SleepFM-2 learned from 282,511 overnight recordings and transferred across diseases, sleep events, wearable sensors and subjective reports.",
      "body_text": "SleepFM-2 was developed and evaluated on 282,511 polysomnography recordings from 26 cohorts, with 235,865 recordings used for pretraining and more than two million hours of physiology in total. The model combines signals from the brain, heart, muscles and respiratory system. With age, sex and BMI, its representation met a prespecified criterion for 215 later-recorded EHR phenotypes in two held-out cohorts; for 155, the sleep representation added reproducible information beyond demographics. The frozen encoder also transferred to expert-scored sleep events, wakeful EEG, headband and in-ear EEG, wrist photoplethysmography and accelerometry. These are retrospective model results across the reported cohorts, not a clinical diagnosis or proof of benefit for an individual.",
      "why_it_matters": "SleepFM-2 learned from 282,511 overnight recordings and transferred across diseases, sleep events, wearable sensors and subjective reports.",
      "limitations": [
        "These are retrospective model results across the reported cohorts, not a clinical diagnosis or proof of benefit for an individual."
      ],
      "importance": 10,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-002/two-million-hours-of-sleep-became-one-health-model",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-002.json",
      "first_published_at": "2026-09-09T09:00:00.000-04:00",
      "modified_at": "2026-09-09T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-09-002-001",
          "text": "SleepFM-2 learned from 282,511 overnight recordings and transferred across diseases, sleep events, wearable sensors and subjective reports.",
          "source_ids": [
            "source-2026-09-09-002"
          ],
          "qualification": "These are retrospective model results across the reported cohorts, not a clinical diagnosis or proof of benefit for an individual."
        }
      ],
      "source_ids": [
        "source-2026-09-09-002"
      ],
      "tags": [
        "sleep physiology",
        "foundation models",
        "wearables"
      ],
      "image_url": "https://themachinepress.com/issues/2026-09-09/feature-sleepfm2.png",
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-003",
      "source_story_id": "tmp-story-setwise-runtime-assurance",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 3,
      "story_type": "dispatch",
      "section": "safety-security",
      "editorial_classification": "editorial",
      "headline": "Retrying a “Safe” Proposal Multiplied Its Risk",
      "slug": "retrying-a-safe-proposal-multiplied-its-risk",
      "dek": "Per-candidate certification can inflate false admission under best-of-k selection; the paper derives a setwise condition that survives generator replacement.",
      "summary": "Per-candidate certification can inflate false admission under best-of-k selection; the paper derives a setwise condition that survives generator replacement.",
      "body_text": "The analysis asks when a runtime safety gate can remain trustworthy even if the learned policy or language-model planner behind it changes. Certifying each candidate separately does not compose: with k retries, a false-admission rate alpha can grow to one minus one-minus-alpha raised to k. The authors prove that simultaneous setwise soundness is necessary and sufficient for generator-independent admission soundness, provided there is a design-time certificate and no bypass path. A second result formalizes what partial observation makes impossible when different hidden states demand different safe actions. The contribution is a theoretical contract and risk ledger, not evidence from a deployed controller.",
      "why_it_matters": "Per-candidate certification can inflate false admission under best-of-k selection; the paper derives a setwise condition that survives generator replacement.",
      "limitations": [
        "The analysis asks when a runtime safety gate can remain trustworthy even if the learned policy or language-model planner behind it changes.",
        "Certifying each candidate separately does not compose: with k retries, a false-admission rate alpha can grow to one minus one-minus-alpha raised to k.",
        "The contribution is a theoretical contract and risk ledger, not evidence from a deployed controller."
      ],
      "importance": 9,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-003/retrying-a-safe-proposal-multiplied-its-risk",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-003.json",
      "first_published_at": "2026-09-09T09:00:00.000-04:00",
      "modified_at": "2026-09-09T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-09-003-001",
          "text": "Per-candidate certification can inflate false admission under best-of-k selection; the paper derives a setwise condition that survives generator replacement.",
          "source_ids": [
            "source-2026-09-09-003"
          ],
          "qualification": "The analysis asks when a runtime safety gate can remain trustworthy even if the learned policy or language-model planner behind it changes."
        }
      ],
      "source_ids": [
        "source-2026-09-09-003"
      ],
      "tags": [
        "runtime assurance",
        "partial observation",
        "safety gates"
      ],
      "image_url": "https://themachinepress.com/issues/2026-09-09/runtime-assurance-laptop-file.webp",
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-004",
      "source_story_id": "tmp-story-darebench-agent-workload-matrix",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 4,
      "story_type": "dispatch",
      "section": "frontier-models",
      "editorial_classification": "editorial",
      "headline": "No Agent Model Won Every Kind of Work",
      "slug": "no-agent-model-won-every-kind-of-work",
      "dek": "DAREBench placed 233 tasks into six workload groups and compared 35 commercial and open models over 7,587 runs.",
      "summary": "DAREBench placed 233 tasks into six workload groups and compared 35 commercial and open models over 7,587 runs.",
      "body_text": "DAREBench adapts tasks from 22 source benchmarks into a two-by-three matrix defined by input modality and execution form. All tasks run in a shared environment with contract-based scoring and evidence audits. Across 23 commercial API models and 12 locally deployed open-weight models, no system dominated all groups. Text and multimodal workloads produced different accuracy-cost trade-offs, while local models were competitive in several groups but trailed frontier commercial systems overall. The benchmark argues that deployment choices should follow workload profiles rather than a single aggregate score; its conclusions remain tied to the selected tasks, environment and reference cost assumptions.",
      "why_it_matters": "DAREBench placed 233 tasks into six workload groups and compared 35 commercial and open models over 7,587 runs.",
      "limitations": [
        "The benchmark argues that deployment choices should follow workload profiles rather than a single aggregate score; its conclusions remain tied to the selected tasks, environment and reference cost assumptions."
      ],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-004/no-agent-model-won-every-kind-of-work",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-004.json",
      "first_published_at": "2026-09-09T09:00:00.000-04:00",
      "modified_at": "2026-09-09T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-09-004-001",
          "text": "DAREBench placed 233 tasks into six workload groups and compared 35 commercial and open models over 7,587 runs.",
          "source_ids": [
            "source-2026-09-09-004"
          ],
          "qualification": "The benchmark argues that deployment choices should follow workload profiles rather than a single aggregate score; its conclusions remain tied to the selected tasks, environment and reference cost assumptions."
        }
      ],
      "source_ids": [
        "source-2026-09-09-004"
      ],
      "tags": [
        "agent benchmarks",
        "model selection",
        "deployment"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-005",
      "source_story_id": "tmp-story-agent-explanations-execution-traces",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 5,
      "story_type": "dispatch",
      "section": "research",
      "editorial_classification": "editorial",
      "headline": "The Explanation Had to Follow the Agent’s Footsteps",
      "slug": "the-explanation-had-to-follow-the-agent-s-footsteps",
      "dek": "A post-hoc framework turns observable execution traces into reports designed to flag unsupported claims, unjustified actions and evidence gaps.",
      "summary": "A post-hoc framework turns observable execution traces into reports designed to flag unsupported claims, unjustified actions and evidence gaps.",
      "body_text": "Traditional explainability methods focus on model outputs or feature influence, while tool-using agents leave a sequence of observable decisions. This framework structures a long execution trace, then produces a natural-language explanation grounded in that record. Human and automated evaluations across multiple benchmarks and agent architectures report better trace faithfulness and stronger identification of unsupported claims, unjustified actions and evidence gaps than naive language-model explanations. Because the method sees behavior rather than hidden reasoning, it can audit what happened without claiming access to an agent’s private internal state.",
      "why_it_matters": "A post-hoc framework turns observable execution traces into reports designed to flag unsupported claims, unjustified actions and evidence gaps.",
      "limitations": [],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-005/the-explanation-had-to-follow-the-agent-s-footsteps",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-005.json",
      "first_published_at": "2026-09-09T09:00:00.000-04:00",
      "modified_at": "2026-09-09T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-09-005-001",
          "text": "A post-hoc framework turns observable execution traces into reports designed to flag unsupported claims, unjustified actions and evidence gaps.",
          "source_ids": [
            "source-2026-09-09-005"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-09-09-005"
      ],
      "tags": [
        "agent traces",
        "explainability",
        "auditing"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-006",
      "source_story_id": "tmp-story-layerroute-vla-representation-routing",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 6,
      "story_type": "dispatch",
      "section": "robotics",
      "editorial_classification": "editorial",
      "headline": "The Robot Chose Which Layer to Remember Mid-Action",
      "slug": "the-robot-chose-which-layer-to-remember-mid-action",
      "dek": "LayerRoute dynamically mixes visual-language layers and rereads earlier action states, adding as little as 0.31 percent in one tested policy.",
      "summary": "LayerRoute dynamically mixes visual-language layers and rereads earlier action states, adding as little as 0.31 percent in one tested policy.",
      "body_text": "Vision-language-action policies commonly expose fixed visual-language layers to each action layer and leave intermediate action states implicit. LayerRoute adds two interfaces: a router that mixes cached visual-language representations according to the current action state, and a reread path for earlier action representations. Across simulation and real-robot benchmarks, the authors report consistent gains for two base policies, including up to 7.2 points on LIBERO Long with 0.31 percent and 3.87 percent additional parameters in the respective systems. The results support adaptive representation access, but do not establish a universal routing recipe for every robot or task.",
      "why_it_matters": "LayerRoute dynamically mixes visual-language layers and rereads earlier action states, adding as little as 0.31 percent in one tested policy.",
      "limitations": [
        "LayerRoute adds two interfaces: a router that mixes cached visual-language representations according to the current action state, and a reread path for earlier action representations.",
        "The results support adaptive representation access, but do not establish a universal routing recipe for every robot or task."
      ],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-006/the-robot-chose-which-layer-to-remember-mid-action",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-006.json",
      "first_published_at": "2026-09-09T09:00:00.000-04:00",
      "modified_at": "2026-09-09T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-09-006-001",
          "text": "LayerRoute dynamically mixes visual-language layers and rereads earlier action states, adding as little as 0.31 percent in one tested policy.",
          "source_ids": [
            "source-2026-09-09-006"
          ],
          "qualification": "LayerRoute adds two interfaces: a router that mixes cached visual-language representations according to the current action state, and a reread path for earlier action representations."
        }
      ],
      "source_ids": [
        "source-2026-09-09-006"
      ],
      "tags": [
        "vision-language-action",
        "robot manipulation",
        "representation routing"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-007",
      "source_story_id": "tmp-story-portable-agent-execution-substrates",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 7,
      "story_type": "dispatch",
      "section": "developer-tools",
      "editorial_classification": "editorial",
      "headline": "One Agent Workflow Ran Three Different Ways",
      "slug": "one-agent-workflow-ran-three-different-ways",
      "dek": "A production platform compiled the same typed graph to streaming, durable orchestration and batch execution without workflow-code changes.",
      "summary": "A production platform compiled the same typed graph to streaming, durable orchestration and batch execution without workflow-code changes.",
      "body_text": "The reported platform emerged from Amazon’s Rufus assistant, where real-time serving, background tasks and high-volume evaluation normally require different runtimes. Developers define one typed dataflow graph, which is then bound to in-process streaming, durable AWS SWF orchestration or distributed Apache Flink processing. Language-model calls become suspendable nodes whose delivery, retry and batching semantics follow the selected substrate. Dozens of production configurations across five orchestration patterns showed no detectable output-quality difference across bindings, while batch execution reduced inference cost in line with published batch pricing. The evidence is an engineering report from one production ecosystem rather than a cross-platform standard.",
      "why_it_matters": "A production platform compiled the same typed graph to streaming, durable orchestration and batch execution without workflow-code changes.",
      "limitations": [],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-007/one-agent-workflow-ran-three-different-ways",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-007.json",
      "first_published_at": "2026-09-09T09:00:00.000-04:00",
      "modified_at": "2026-09-09T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-09-007-001",
          "text": "A production platform compiled the same typed graph to streaming, durable orchestration and batch execution without workflow-code changes.",
          "source_ids": [
            "source-2026-09-09-007"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-09-09-007"
      ],
      "tags": [
        "agent workflows",
        "runtime portability",
        "batch inference"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-008",
      "source_story_id": "tmp-story-scirigor-evidence-chain-audit",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 8,
      "story_type": "dispatch",
      "section": "research",
      "editorial_classification": "editorial",
      "headline": "Correct-Looking Claims Survived Broken Evidence",
      "slug": "correct-looking-claims-survived-broken-evidence",
      "dek": "In SciRIGOR, claims agreed with faithful and unfaithful results at almost the same rate, while strict end-to-end evidence success stayed below 18 percent.",
      "summary": "In SciRIGOR, claims agreed with faithful and unfaithful results at almost the same rate, while strict end-to-end evidence success stayed below 18 percent.",
      "body_text": "SciRIGOR evaluates scientific coding agents as linked chains from executable analysis through results and figures to claims. Its 100 cases span six domains and 17 subfields, with typed evidence graphs that distinguish artifact fidelity from the validity of each supporting relation. Across 11 agent-model configurations, claims agreed with faithful results 91.8 percent of the time and with unfaithful results 91.0 percent of the time. No system exceeded 62.6 percent on the soft evidence-chain score or 18 percent on strict whole-chain success. Internal coherence therefore did not establish scientific correctness in this benchmark.",
      "why_it_matters": "In SciRIGOR, claims agreed with faithful and unfaithful results at almost the same rate, while strict end-to-end evidence success stayed below 18 percent.",
      "limitations": [
        "Internal coherence therefore did not establish scientific correctness in this benchmark."
      ],
      "importance": 9,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-008/correct-looking-claims-survived-broken-evidence",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-008.json",
      "first_published_at": "2026-09-09T09:00:00.000-04:00",
      "modified_at": "2026-09-09T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-09-008-001",
          "text": "In SciRIGOR, claims agreed with faithful and unfaithful results at almost the same rate, while strict end-to-end evidence success stayed below 18 percent.",
          "source_ids": [
            "source-2026-09-09-008"
          ],
          "qualification": "Internal coherence therefore did not establish scientific correctness in this benchmark."
        }
      ],
      "source_ids": [
        "source-2026-09-09-008"
      ],
      "tags": [
        "scientific agents",
        "evidence chains",
        "evaluation"
      ],
      "image_url": "https://themachinepress.com/issues/2026-09-09/scirigor-code-file.webp",
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-009",
      "source_story_id": "tmp-story-normviz-global-visual-norms",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 9,
      "story_type": "dispatch",
      "section": "consumer-products",
      "editorial_classification": "editorial",
      "headline": "The Best Vision Model Read Fewer Than Three in Ten Cultural Pairs",
      "slug": "the-best-vision-model-read-fewer-than-three-in-ten-cultural-pairs",
      "dek": "NormViz-Bench changed one culturally relevant behavior at a time across 6,536 images from 16 countries.",
      "summary": "NormViz-Bench changed one culturally relevant behavior at a time across 6,536 images from 16 countries.",
      "body_text": "NormViz-Bench contains 3,268 human-validated contrastive image pairs spanning 16 countries. Each pair differs only in a behavior that changes whether the scene conforms to, violates or is irrelevant to a local social norm, and pair-level scoring requires both images to be correct. Gemini 3.0 Flash and Qwen2.5-VL-7B reached 26.6 and 21.6 percent pair accuracy, respectively. Fine-tuning smaller Qwen3-VL models on 64,000 explanation-linked images raised relative accuracy, but absolute performance remained below 30 percent. The benchmark measures its curated countries, behaviors and labels; it does not reduce culture to a single universal rulebook.",
      "why_it_matters": "NormViz-Bench changed one culturally relevant behavior at a time across 6,536 images from 16 countries.",
      "limitations": [
        "Each pair differs only in a behavior that changes whether the scene conforms to, violates or is irrelevant to a local social norm, and pair-level scoring requires both images to be correct.",
        "The benchmark measures its curated countries, behaviors and labels; it does not reduce culture to a single universal rulebook."
      ],
      "importance": 9,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-009/the-best-vision-model-read-fewer-than-three-in-ten-cultural-pairs",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-009.json",
      "first_published_at": "2026-09-09T09:00:00.000-04:00",
      "modified_at": "2026-09-09T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-09-009-001",
          "text": "NormViz-Bench changed one culturally relevant behavior at a time across 6,536 images from 16 countries.",
          "source_ids": [
            "source-2026-09-09-009"
          ],
          "qualification": "Each pair differs only in a behavior that changes whether the scene conforms to, violates or is irrelevant to a local social norm, and pair-level scoring requires both images to be correct."
        }
      ],
      "source_ids": [
        "source-2026-09-09-009"
      ],
      "tags": [
        "multimodal AI",
        "culture",
        "benchmarking"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-010",
      "source_story_id": "tmp-story-thoughtmed-social-media-medical-vqa",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 10,
      "story_type": "dispatch",
      "section": "research",
      "editorial_classification": "editorial",
      "headline": "A Million Medical Questions Came From Clinicians’ Shared Images",
      "slug": "a-million-medical-questions-came-from-clinicians-shared-images",
      "dek": "ThoughtMed-1M pairs de-identified images with verified expert commentary; its trained model reached 85.4 percent macro accuracy across 42 benchmarks.",
      "summary": "ThoughtMed-1M pairs de-identified images with verified expert commentary; its trained model reached 85.4 percent macro accuracy across 42 benchmarks.",
      "body_text": "The authors built a pipeline around de-identified medical images and commentaries shared on clinician-oriented social media, using a language model plus clinician-in-the-loop verification to create more than one million long-form visual question-answer pairs. A foundation model trained on that set, FOLTMed, achieved 85.4 percent macro accuracy across 42 medical VQA benchmarks and outscored comparison systems by three to five points on reported factuality and similarity measures. The work presents a scalable research dataset and benchmark result, not clinical approval, and its provenance and de-identification controls remain central to responsible reuse.",
      "why_it_matters": "ThoughtMed-1M pairs de-identified images with verified expert commentary; its trained model reached 85.4 percent macro accuracy across 42 benchmarks.",
      "limitations": [
        "The work presents a scalable research dataset and benchmark result, not clinical approval, and its provenance and de-identification controls remain central to responsible reuse."
      ],
      "importance": 9,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-010/a-million-medical-questions-came-from-clinicians-shared-images",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-010.json",
      "first_published_at": "2026-09-09T09:00:00.000-04:00",
      "modified_at": "2026-09-09T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-09-010-001",
          "text": "ThoughtMed-1M pairs de-identified images with verified expert commentary; its trained model reached 85.4 percent macro accuracy across 42 benchmarks.",
          "source_ids": [
            "source-2026-09-09-010"
          ],
          "qualification": "The work presents a scalable research dataset and benchmark result, not clinical approval, and its provenance and de-identification controls remain central to responsible reuse."
        }
      ],
      "source_ids": [
        "source-2026-09-09-010"
      ],
      "tags": [
        "medical imaging",
        "visual question answering",
        "clinical datasets"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-011",
      "source_story_id": "tmp-story-ibrain-surface-to-spikes",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 11,
      "story_type": "dispatch",
      "section": "research",
      "editorial_classification": "editorial",
      "headline": "Seven Thousand Hours Joined Brain-Surface Signals to Spikes",
      "slug": "seven-thousand-hours-joined-brain-surface-signals-to-spikes",
      "dek": "iBrain jointly pretrained on intracranial EEG and intracortical spiking with signal-specific encoders and one shared temporal backbone.",
      "summary": "iBrain jointly pretrained on intracranial EEG and intracortical spiking with signal-specific encoders and one shared temporal backbone.",
      "body_text": "Most neural foundation models specialize in one recording type. iBrain instead uses separate encoders for intracranial EEG and intracortical spike trains, followed by a shared spatiotemporal transformer trained with masked reconstruction and channel-view alignment. The pretraining corpus contains more than 7,000 hours of heterogeneous invasive recordings. Across the reported benchmarks, the joint model outperformed single-signal pretraining baselines and transferred with improved data efficiency across recording settings. These are research benchmarks on invasive recordings, not evidence that the model can read thoughts or support unsupervised clinical decisions.",
      "why_it_matters": "iBrain jointly pretrained on intracranial EEG and intracortical spiking with signal-specific encoders and one shared temporal backbone.",
      "limitations": [
        "These are research benchmarks on invasive recordings, not evidence that the model can read thoughts or support unsupervised clinical decisions."
      ],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-011/seven-thousand-hours-joined-brain-surface-signals-to-spikes",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-011.json",
      "first_published_at": "2026-09-09T09:00:00.000-04:00",
      "modified_at": "2026-09-09T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-09-011-001",
          "text": "iBrain jointly pretrained on intracranial EEG and intracortical spiking with signal-specific encoders and one shared temporal backbone.",
          "source_ids": [
            "source-2026-09-09-011"
          ],
          "qualification": "These are research benchmarks on invasive recordings, not evidence that the model can read thoughts or support unsupervised clinical decisions."
        }
      ],
      "source_ids": [
        "source-2026-09-09-011"
      ],
      "tags": [
        "neural recordings",
        "foundation models",
        "intracranial EEG"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-012",
      "source_story_id": "tmp-story-physics-constrained-vehicle-platoons",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 12,
      "story_type": "dispatch",
      "section": "infrastructure",
      "editorial_classification": "editorial",
      "headline": "The Car Forecast Stayed Accurate Without Amplifying the Shockwave",
      "slug": "the-car-forecast-stayed-accurate-without-amplifying-the-shockwave",
      "dek": "A platoon model added learned propagation delays and string-stability losses, keeping unstable windows to 0.65 percent on one five-car test.",
      "summary": "A platoon model added learned propagation delays and string-stability losses, keeping unstable windows to 0.65 percent on one five-car test.",
      "body_text": "SSP-DMGTimeNet predicts several vehicles together while penalizing disturbance amplification through a platoon. Its attention mechanism learns response delays between adjacent cars and accumulates them downstream; time- and frequency-domain losses target string stability for neighboring vehicles and longer sub-platoons. On the HighD ground-truth excitation subset, the reported five-vehicle unstable-window rate was 0.65 percent and maximum head-to-tail amplification was 0.898. Zero-shot tests on NGSIM US-101 and I-80 produced unstable-window rates of 3.90 and 4.10 percent. The figures are dataset results, not a road-safety validation.",
      "why_it_matters": "A platoon model added learned propagation delays and string-stability losses, keeping unstable windows to 0.65 percent on one five-car test.",
      "limitations": [
        "The figures are dataset results, not a road-safety validation."
      ],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-012/the-car-forecast-stayed-accurate-without-amplifying-the-shockwave",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-012.json",
      "first_published_at": "2026-09-09T09:00:00.000-04:00",
      "modified_at": "2026-09-09T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-09-012-001",
          "text": "A platoon model added learned propagation delays and string-stability losses, keeping unstable windows to 0.65 percent on one five-car test.",
          "source_ids": [
            "source-2026-09-09-012"
          ],
          "qualification": "The figures are dataset results, not a road-safety validation."
        }
      ],
      "source_ids": [
        "source-2026-09-09-012"
      ],
      "tags": [
        "vehicle platoons",
        "trajectory prediction",
        "physics constraints"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-013",
      "source_story_id": "tmp-story-fixed-shallow-classical-shadow-analyzer",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 13,
      "story_type": "dispatch",
      "section": "research",
      "editorial_classification": "editorial",
      "headline": "One Fixed Circuit Replaced Exponentially Many Measurement Settings",
      "slug": "one-fixed-circuit-replaced-exponentially-many-measurement-settings",
      "dek": "A shallow analyzer turns Born outcomes into reusable classical-shadow labels while retaining the minimum number of outcomes for complete reconstruction.",
      "summary": "A shallow analyzer turns Born outcomes into reusable classical-shadow labels while retaining the minimum number of outcomes for complete reconstruction.",
      "body_text": "Classical shadows usually randomize among measurement settings, adding control and calibration overhead beyond circuit depth. This construction couples the unknown n-qubit system to a freshly prepared n-qubit fiducial register and performs parallel Bell readout. One fixed setting then replaces the 3-to-the-n local-Pauli settings used for complete reconstruction while retaining d-squared outcomes for d equals 2 to the n. The fiducial preparation uses n minus one arbitrary two-qubit gates and logarithmic depth with all-to-all connectivity; the unknown system sees one parallel entangling layer. The result is theoretical trade-off analysis, not a demonstration on a specific quantum processor.",
      "why_it_matters": "A shallow analyzer turns Born outcomes into reusable classical-shadow labels while retaining the minimum number of outcomes for complete reconstruction.",
      "limitations": [
        "The result is theoretical trade-off analysis, not a demonstration on a specific quantum processor."
      ],
      "importance": 9,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-013/one-fixed-circuit-replaced-exponentially-many-measurement-settings",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-013.json",
      "first_published_at": "2026-09-09T09:00:00.000-04:00",
      "modified_at": "2026-09-09T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-09-013-001",
          "text": "A shallow analyzer turns Born outcomes into reusable classical-shadow labels while retaining the minimum number of outcomes for complete reconstruction.",
          "source_ids": [
            "source-2026-09-09-013"
          ],
          "qualification": "The result is theoretical trade-off analysis, not a demonstration on a specific quantum processor."
        }
      ],
      "source_ids": [
        "source-2026-09-09-013"
      ],
      "tags": [
        "classical shadows",
        "quantum measurement",
        "shallow circuits"
      ],
      "image_url": "https://themachinepress.com/issues/2026-09-09/quantum-shadow-chip-file.webp",
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-014",
      "source_story_id": "tmp-story-quantum-decoder-conformance-contract",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 14,
      "story_type": "dispatch",
      "section": "developer-tools",
      "editorial_classification": "editorial",
      "headline": "Quantum Decoder Libraries Declared Four of 54 Capabilities",
      "slug": "quantum-decoder-libraries-declared-four-of-54-capabilities",
      "dek": "An oracle-free contract caught bookkeeping-sensitive behavior that ordinary logical-error counting can miss.",
      "summary": "An oracle-free contract caught bookkeeping-sensitive behavior that ordinary logical-error counting can miss.",
      "body_text": "The proposed conformance suite checks whether a decoder’s correction explains the syndrome in the caller’s index space and whether equivalent presentations can force a correction heavier than another known feasible one. Verdicts are gated on each library’s own declarations. Across nine configurations from five public libraries, documentation answered four of 54 capability questions, and none directly declared bounded-distance correctness even though its hypotheses held in 62.1 percent of cases. One solver returned a correction 26 percent heavier after only the numbering changed. All 639 certificates ship with a standalone verifier, making the reported contradictions independently re-derivable.",
      "why_it_matters": "An oracle-free contract caught bookkeeping-sensitive behavior that ordinary logical-error counting can miss.",
      "limitations": [
        "One solver returned a correction 26 percent heavier after only the numbering changed."
      ],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-014/quantum-decoder-libraries-declared-four-of-54-capabilities",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-014.json",
      "first_published_at": "2026-09-09T09:00:00.000-04:00",
      "modified_at": "2026-09-09T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-09-014-001",
          "text": "An oracle-free contract caught bookkeeping-sensitive behavior that ordinary logical-error counting can miss.",
          "source_ids": [
            "source-2026-09-09-014"
          ],
          "qualification": "One solver returned a correction 26 percent heavier after only the numbering changed."
        }
      ],
      "source_ids": [
        "source-2026-09-09-014"
      ],
      "tags": [
        "quantum error correction",
        "software testing",
        "decoder libraries"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-026",
      "source_story_id": "tmp-story-moonlit-satellite-observations",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 15,
      "story_type": "dispatch",
      "section": "research",
      "editorial_classification": "editorial",
      "headline": "Moonlight Kept the Space Stations Visible Through Midnight",
      "slug": "moonlight-kept-the-space-stations-visible-through-midnight",
      "dek": "A 0.6-meter telescope recorded 147 night detections and a model found lunar Earthshine can boost nadir-facing radiance by at least 100 times over starlight.",
      "summary": "A 0.6-meter telescope recorded 147 night detections and a model found lunar Earthshine can boost nadir-facing radiance by at least 100 times over starlight.",
      "body_text": "Optical tracking usually depends on sunlit satellite passes near twilight. The authors add moonlight and lunar Earthshine to sunlight and Earthshine in a satellite-brightness model, then compare it with what they describe as the first quantitative night observations of the ISS and Chinese Space Station illuminated only by lunar sources. The 147 ISS detections had median brightness V equals 12.02 plus or minus 0.17, about 13.1 magnitudes fainter than daylight. Their model suggests large spacecraft can remain within reach of modest telescopes for much of the night and up to roughly 11 nights per lunar month at local midnight.",
      "why_it_matters": "A 0.6-meter telescope recorded 147 night detections and a model found lunar Earthshine can boost nadir-facing radiance by at least 100 times over starlight.",
      "limitations": [
        "The authors add moonlight and lunar Earthshine to sunlight and Earthshine in a satellite-brightness model, then compare it with what they describe as the first quantitative night observations of the ISS and Chinese Space Station illuminated only by lunar sources.",
        "Their model suggests large spacecraft can remain within reach of modest telescopes for much of the night and up to roughly 11 nights per lunar month at local midnight."
      ],
      "importance": 9,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-026/moonlight-kept-the-space-stations-visible-through-midnight",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-026.json",
      "first_published_at": "2026-09-09T09:00:00.000-04:00",
      "modified_at": "2026-09-09T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-09-026-001",
          "text": "A 0.6-meter telescope recorded 147 night detections and a model found lunar Earthshine can boost nadir-facing radiance by at least 100 times over starlight.",
          "source_ids": [
            "source-2026-09-09-015"
          ],
          "qualification": "The authors add moonlight and lunar Earthshine to sunlight and Earthshine in a satellite-brightness model, then compare it with what they describe as the first quantitative night observations of the ISS and Chinese Space Station illuminated only by lunar sources."
        }
      ],
      "source_ids": [
        "source-2026-09-09-015"
      ],
      "tags": [
        "satellite tracking",
        "moonlight",
        "space domain awareness"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-027",
      "source_story_id": "tmp-story-astrospeclm-grounded-spectra-front-rail",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 16,
      "story_type": "dispatch",
      "section": "research",
      "editorial_classification": "editorial",
      "headline": "A Language Model Learned to Point Back to the Spectrum",
      "slug": "a-language-model-learned-to-point-back-to-the-spectrum",
      "dek": "AstroSpecLM turns DESI spectra into fact-grounded conversations, combining classification and redshift prediction with feature-linked explanations.",
      "summary": "AstroSpecLM turns DESI spectra into fact-grounded conversations, combining classification and redshift prediction with feature-linked explanations.",
      "body_text": "AstroSpecLM connects one-dimensional DESI spectra with Qwen3-4B. Instead of generating training conversations directly from templates or raw catalog fields, the pipeline first distills each spectrum into a compact set of catalog- and spectrum-derived facts. Those facts become references for instruction-following examples. The resulting model was competitive with specialist supervised baselines on classification and redshift estimation while producing explanations that cited specific spectral features. The paper establishes feasibility on the reported data; natural-language fluency does not independently validate every astronomical interpretation.",
      "why_it_matters": "AstroSpecLM turns DESI spectra into fact-grounded conversations, combining classification and redshift prediction with feature-linked explanations.",
      "limitations": [
        "The paper establishes feasibility on the reported data; natural-language fluency does not independently validate every astronomical interpretation."
      ],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-027/a-language-model-learned-to-point-back-to-the-spectrum",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-027.json",
      "first_published_at": "2026-09-09T09:00:00.000-04:00",
      "modified_at": "2026-09-09T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-09-027-001",
          "text": "AstroSpecLM turns DESI spectra into fact-grounded conversations, combining classification and redshift prediction with feature-linked explanations.",
          "source_ids": [
            "source-2026-09-09-016"
          ],
          "qualification": "The paper establishes feasibility on the reported data; natural-language fluency does not independently validate every astronomical interpretation."
        }
      ],
      "source_ids": [
        "source-2026-09-09-016"
      ],
      "tags": [
        "astronomical spectra",
        "DESI",
        "grounded explanations"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-015",
      "source_story_id": "tmp-sidebar-hierarchical-rag-consistency",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 17,
      "story_type": "ticker",
      "section": "developer-tools",
      "editorial_classification": "editorial",
      "headline": "A Correct Answer Could Still Hide a Contradictory Source",
      "slug": "a-correct-answer-could-still-hide-a-contradictory-source",
      "dek": "A three-level audit separates conflicts in the corpus, retrieved context and generated answer rather than scoring only the final response.",
      "summary": "A three-level audit separates conflicts in the corpus, retrieved context and generated answer rather than scoring only the final response.",
      "body_text": "Across 100 controlled query-corpus cases in five domains, the framework found that retrieval similarity, corpus quality and answer consistency can move in different directions. It makes supporting and contradictory source-linked facts inspectable but does not certify external truth.",
      "why_it_matters": "A three-level audit separates conflicts in the corpus, retrieved context and generated answer rather than scoring only the final response.",
      "limitations": [
        "It makes supporting and contradictory source-linked facts inspectable but does not certify external truth."
      ],
      "importance": 7,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-015/a-correct-answer-could-still-hide-a-contradictory-source",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-015.json",
      "first_published_at": "2026-09-09T09:00:00.000-04:00",
      "modified_at": "2026-09-09T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-09-015-001",
          "text": "A three-level audit separates conflicts in the corpus, retrieved context and generated answer rather than scoring only the final response.",
          "source_ids": [
            "source-2026-09-09-017"
          ],
          "qualification": "It makes supporting and contradictory source-linked facts inspectable but does not certify external truth."
        }
      ],
      "source_ids": [
        "source-2026-09-09-017"
      ],
      "tags": [
        "RAG",
        "contradictions",
        "auditing"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-016",
      "source_story_id": "tmp-sidebar-review-value-exposure-budget",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 18,
      "story_type": "ticker",
      "section": "safety-security",
      "editorial_classification": "editorial",
      "headline": "The Riskiest Answer Was Not Always the Best One to Review",
      "slug": "the-riskiest-answer-was-not-always-the-best-one-to-review",
      "dek": "At a 20 percent review budget, ranking by repair value cut residual wrong-answer exposure from 0.881 to 0.716 in a 720-item stress test.",
      "summary": "At a 20 percent review budget, ranking by repair value cut residual wrong-answer exposure from 0.881 to 0.716 in a 720-item stress test.",
      "body_text": "The proposed review queue combines estimated wrongness with intervention affordance, impact and cost. On TAT-QA and SciFact items, it barely changed wrong-answer exposure before repair but substantially reduced exposure after deterministic benchmark-supported fixes.",
      "why_it_matters": "At a 20 percent review budget, ranking by repair value cut residual wrong-answer exposure from 0.881 to 0.716 in a 720-item stress test.",
      "limitations": [],
      "importance": 7,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-016/the-riskiest-answer-was-not-always-the-best-one-to-review",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-016.json",
      "first_published_at": "2026-09-09T09:00:00.000-04:00",
      "modified_at": "2026-09-09T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-09-016-001",
          "text": "At a 20 percent review budget, ranking by repair value cut residual wrong-answer exposure from 0.881 to 0.716 in a 720-item stress test.",
          "source_ids": [
            "source-2026-09-09-018"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-09-09-018"
      ],
      "tags": [
        "human review",
        "risk prioritization",
        "repairability"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-017",
      "source_story_id": "tmp-sidebar-occupancy-world-model-filtering",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 19,
      "story_type": "ticker",
      "section": "robotics",
      "editorial_classification": "editorial",
      "headline": "A More Accurate 3D Map Did Not Always Choose a Better View",
      "slug": "a-more-accurate-3d-map-did-not-always-choose-a-better-view",
      "dek": "Correcting false positive or false negative occupancy alone failed to consistently improve final coverage under a fixed active-mapping planner.",
      "summary": "Correcting false positive or false negative occupancy alone failed to consistently improve final coverage under a fixed active-mapping planner.",
      "body_text": "The diagnosis separates geometric completion from downstream planning. A preliminary filter suppresses repeatedly unsupported occupancy while preserving predictions in unexplored space, redirecting some viewpoints toward reachable surfaces.",
      "why_it_matters": "Correcting false positive or false negative occupancy alone failed to consistently improve final coverage under a fixed active-mapping planner.",
      "limitations": [],
      "importance": 7,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-017/a-more-accurate-3d-map-did-not-always-choose-a-better-view",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-017.json",
      "first_published_at": "2026-09-09T09:00:00.000-04:00",
      "modified_at": "2026-09-09T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-09-017-001",
          "text": "Correcting false positive or false negative occupancy alone failed to consistently improve final coverage under a fixed active-mapping planner.",
          "source_ids": [
            "source-2026-09-09-019"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-09-09-019"
      ],
      "tags": [
        "active mapping",
        "occupancy models",
        "view planning"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-018",
      "source_story_id": "tmp-sidebar-memobench-robot-memory-process",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 20,
      "story_type": "ticker",
      "section": "robotics",
      "editorial_classification": "editorial",
      "headline": "The Robot Remembered the Scene but Mishandled the Update",
      "slug": "the-robot-remembered-the-scene-but-mishandled-the-update",
      "dek": "MEMOBench labels storage, update and compression separately across 4,200 checkpoints; its strongest memory baseline averaged 31.9 percent success.",
      "summary": "MEMOBench labels storage, update and compression separately across 4,200 checkpoints; its strongest memory baseline averaged 31.9 percent success.",
      "body_text": "The suite contains 30 history-dependent manipulation tasks, 1,500 expert demonstrations and checkpoints from 84 templates. High storage scores often coexisted with weak updating and compression, separating memory errors from manipulation failure.",
      "why_it_matters": "MEMOBench labels storage, update and compression separately across 4,200 checkpoints; its strongest memory baseline averaged 31.9 percent success.",
      "limitations": [],
      "importance": 7,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-018/the-robot-remembered-the-scene-but-mishandled-the-update",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-018.json",
      "first_published_at": "2026-09-09T09:00:00.000-04:00",
      "modified_at": "2026-09-09T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-09-018-001",
          "text": "MEMOBench labels storage, update and compression separately across 4,200 checkpoints; its strongest memory baseline averaged 31.9 percent success.",
          "source_ids": [
            "source-2026-09-09-020"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-09-09-020"
      ],
      "tags": [
        "robot memory",
        "manipulation",
        "benchmarks"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-019",
      "source_story_id": "tmp-sidebar-three-band-black-hole-mergers",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 21,
      "story_type": "ticker",
      "section": "research",
      "editorial_classification": "editorial",
      "headline": "Three Frequency Bands Could Catch Two Dozen of the Same Mergers",
      "slug": "three-frequency-bands-could-catch-two-dozen-of-the-same-mergers",
      "dek": "A forecast found the largest three-band yield when decihertz observations begin within about three years of the millihertz mission.",
      "summary": "A forecast found the largest three-band yield when decihertz observations begin within about three years of the millihertz mission.",
      "body_text": "Using GWTC-4-derived populations and explicit mission timing, the study forecasts 23.5 plus 8.7 or minus 6.2 three-band events for one network at signal-to-noise eight. Subthreshold searches could raise yields four- to fivefold; these remain mission and population forecasts.",
      "why_it_matters": "A forecast found the largest three-band yield when decihertz observations begin within about three years of the millihertz mission.",
      "limitations": [
        "Subthreshold searches could raise yields four- to fivefold; these remain mission and population forecasts."
      ],
      "importance": 7,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-019/three-frequency-bands-could-catch-two-dozen-of-the-same-mergers",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-019.json",
      "first_published_at": "2026-09-09T09:00:00.000-04:00",
      "modified_at": "2026-09-09T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-09-09-019-001",
          "text": "A forecast found the largest three-band yield when decihertz observations begin within about three years of the millihertz mission.",
          "source_ids": [
            "source-2026-09-09-021"
          ],
          "qualification": "Subthreshold searches could raise yields four- to fivefold; these remain mission and population forecasts."
        }
      ],
      "source_ids": [
        "source-2026-09-09-021"
      ],
      "tags": [
        "gravitational waves",
        "black holes",
        "multiband astronomy"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-020",
      "source_story_id": "invention-hangprinter",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 22,
      "story_type": "invention_desk",
      "section": "invention-desk",
      "editorial_classification": "editorial",
      "headline": "Hangprinter",
      "slug": "hangprinter",
      "dek": "Suspends a print head from tensioned lines anchored around a room, replacing a rigid gantry with cable geometry so an open RepRap can work across an unusually large build space.",
      "summary": "Suspends a print head from tensioned lines anchored around a room, replacing a rigid gantry with cable geometry so an open RepRap can work across an unusually large build space.",
      "body_text": "Suspends a print head from tensioned lines anchored around a room, replacing a rigid gantry with cable geometry so an open RepRap can work across an unusually large build space.",
      "why_it_matters": "An independent builder is turning an improbable idea into a working project.",
      "limitations": [
        "A Desk Pick is an editorial selection, not a product endorsement."
      ],
      "importance": null,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-020/hangprinter",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-020.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "carried_over",
      "is_carryover": true,
      "carryover_reason": "invention_desk_seven_day_placement",
      "key_claims": [],
      "source_ids": [],
      "tags": [
        "independent-builders",
        "desk-pick",
        "prototype"
      ],
      "image_url": "https://themachinepress.com/issues/2026-09-09/invention-desk/hangprinter.png",
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-021",
      "source_story_id": "invention-precious-plastic",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 23,
      "story_type": "invention_desk",
      "section": "invention-desk",
      "editorial_classification": "editorial",
      "headline": "Precious Plastic",
      "slug": "precious-plastic",
      "dek": "Publishes replicable shredders, presses, workspace plans, and shared know-how so small local teams can sort waste plastic and turn it into reusable flakes and sheet material.",
      "summary": "Publishes replicable shredders, presses, workspace plans, and shared know-how so small local teams can sort waste plastic and turn it into reusable flakes and sheet material.",
      "body_text": "Publishes replicable shredders, presses, workspace plans, and shared know-how so small local teams can sort waste plastic and turn it into reusable flakes and sheet material.",
      "why_it_matters": "An independent builder is turning an improbable idea into a working project.",
      "limitations": [
        "A Desk Pick is an editorial selection, not a product endorsement."
      ],
      "importance": null,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-021/precious-plastic",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-021.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "carried_over",
      "is_carryover": true,
      "carryover_reason": "invention_desk_seven_day_placement",
      "key_claims": [],
      "source_ids": [],
      "tags": [
        "independent-builders",
        "desk-pick",
        "released"
      ],
      "image_url": "https://themachinepress.com/issues/2026-09-09/invention-desk/precious-plastic.png",
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-022",
      "source_story_id": "invention-watchy",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 24,
      "story_type": "invention_desk",
      "section": "invention-desk",
      "editorial_classification": "editorial",
      "headline": "Watchy",
      "slug": "watchy",
      "dek": "Pairs a square e-paper display with an ESP32-S3 and publishes the hardware, software, documentation, and case files so owners can build and program their own watch faces.",
      "summary": "Pairs a square e-paper display with an ESP32-S3 and publishes the hardware, software, documentation, and case files so owners can build and program their own watch faces.",
      "body_text": "Pairs a square e-paper display with an ESP32-S3 and publishes the hardware, software, documentation, and case files so owners can build and program their own watch faces.",
      "why_it_matters": "An independent builder is turning an improbable idea into a working project.",
      "limitations": [
        "A Desk Pick is an editorial selection, not a product endorsement."
      ],
      "importance": null,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-022/watchy",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-022.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "carried_over",
      "is_carryover": true,
      "carryover_reason": "invention_desk_seven_day_placement",
      "key_claims": [],
      "source_ids": [],
      "tags": [
        "independent-builders",
        "desk-pick",
        "released"
      ],
      "image_url": "https://themachinepress.com/issues/2026-09-09/invention-desk/watchy.png",
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-023",
      "source_story_id": "invention-ploopy-classic-2",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 25,
      "story_type": "invention_desk",
      "section": "invention-desk",
      "editorial_classification": "editorial",
      "headline": "Ploopy Classic 2",
      "slug": "ploopy-classic-2",
      "dek": "Turns a desktop trackball into an inspectable kit by publishing its mechanical and electrical design files, assembly documentation, and programmable QMK firmware.",
      "summary": "Turns a desktop trackball into an inspectable kit by publishing its mechanical and electrical design files, assembly documentation, and programmable QMK firmware.",
      "body_text": "Turns a desktop trackball into an inspectable kit by publishing its mechanical and electrical design files, assembly documentation, and programmable QMK firmware.",
      "why_it_matters": "An independent builder is turning an improbable idea into a working project.",
      "limitations": [
        "A Desk Pick is an editorial selection, not a product endorsement."
      ],
      "importance": null,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-023/ploopy-classic-2",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-023.json",
      "first_published_at": "2026-09-06T09:00:00.000-04:00",
      "modified_at": "2026-09-06T09:00:00.000-04:00",
      "content_status": "carried_over",
      "is_carryover": true,
      "carryover_reason": "invention_desk_seven_day_placement",
      "key_claims": [],
      "source_ids": [],
      "tags": [
        "independent-builders",
        "desk-pick",
        "released"
      ],
      "image_url": "https://themachinepress.com/issues/2026-09-09/invention-desk/ploopy-classic-2.png",
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-024",
      "source_story_id": "invention-sponsored-house-example",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 26,
      "story_type": "invention_desk",
      "section": "invention-desk",
      "editorial_classification": "house_example",
      "headline": "The First Paid Slot",
      "slug": "the-first-paid-slot",
      "dek": "A transparent preview of paid placement with one verified link and no claim of endorsement.",
      "summary": "A transparent preview of paid placement with one verified link and no claim of endorsement.",
      "body_text": "A transparent preview of paid placement with one verified link and no claim of endorsement.\n\nHouse example - no advertiser paid. Payment will buy placement, never endorsement.",
      "why_it_matters": "This placement explains how builders can appear in The Invention Desk without purchasing editorial endorsement.",
      "limitations": [
        "House example - no advertiser paid. Payment will buy placement, never endorsement."
      ],
      "importance": null,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-024/the-first-paid-slot",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-024.json",
      "first_published_at": "2026-07-10T09:00:00.000-04:00",
      "modified_at": "2026-07-10T09:00:00.000-04:00",
      "content_status": "carried_over",
      "is_carryover": true,
      "carryover_reason": "invention_desk_seven_day_placement",
      "key_claims": [],
      "source_ids": [],
      "tags": [
        "independent-builders",
        "sponsored-project",
        "open"
      ],
      "image_url": "https://themachinepress.com/issues/2026-09-09/invention-desk/sponsored-project.webp",
      "corrections": []
    },
    {
      "story_id": "mp-2026-09-09-025",
      "source_story_id": "invention-placement-cta",
      "edition_id": "mp-2026-09-09-morning-0062",
      "edition_url": "https://themachinepress.com/edition/2026-09-09",
      "position": 27,
      "story_type": "invention_desk",
      "section": "invention-desk",
      "editorial_classification": "house_example",
      "headline": "Put Your Project on the Desk",
      "slug": "put-your-project-on-the-desk",
      "dek": "One manually reviewed placement stays active for seven days and remains separate from Desk Picks.",
      "summary": "One manually reviewed placement stays active for seven days and remains separate from Desk Picks.",
      "body_text": "One manually reviewed placement stays active for seven days and remains separate from Desk Picks.\n\nManual intake only. Payment buys placement, never endorsement, and every submission is reviewed.",
      "why_it_matters": "This placement explains how builders can appear in The Invention Desk without purchasing editorial endorsement.",
      "limitations": [
        "Manual intake only. Payment buys placement, never endorsement, and every submission is reviewed."
      ],
      "importance": null,
      "canonical_url": "https://themachinepress.com/story/mp-2026-09-09-025/put-your-project-on-the-desk",
      "json_url": "https://themachinepress.com/story/mp-2026-09-09-025.json",
      "first_published_at": "2026-07-10T09:00:00.000-04:00",
      "modified_at": "2026-07-10T09:00:00.000-04:00",
      "content_status": "carried_over",
      "is_carryover": true,
      "carryover_reason": "invention_desk_seven_day_placement",
      "key_claims": [],
      "source_ids": [],
      "tags": [
        "independent-builders",
        "cta",
        "open"
      ],
      "image_url": "https://themachinepress.com/issues/2026-09-09/invention-desk/put-your-project-on-the-desk.webp",
      "corrections": []
    }
  ],
  "sources": [
    {
      "source_id": "source-2026-09-09-001",
      "title": "arXiv preprint 2609.07288",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.07288",
      "canonical_url": "https://arxiv.org/abs/2609.07288",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-07T05:58:52.000-04:00",
      "accessed_at": "2026-09-09T08:30:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-09-001-001"
      ]
    },
    {
      "source_id": "source-2026-09-09-002",
      "title": "arXiv preprint 2609.06849",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.06849",
      "canonical_url": "https://arxiv.org/abs/2609.06849",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-06T17:47:56.000-04:00",
      "accessed_at": "2026-09-09T08:30:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-09-002-001"
      ]
    },
    {
      "source_id": "source-2026-09-09-003",
      "title": "arXiv preprint 2609.06036",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.06036",
      "canonical_url": "https://arxiv.org/abs/2609.06036",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-05T07:40:42.000-04:00",
      "accessed_at": "2026-09-09T08:30:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-09-003-001"
      ]
    },
    {
      "source_id": "source-2026-09-09-004",
      "title": "arXiv preprint 2609.06059",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.06059",
      "canonical_url": "https://arxiv.org/abs/2609.06059",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-05T08:36:59.000-04:00",
      "accessed_at": "2026-09-09T08:30:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-09-004-001"
      ]
    },
    {
      "source_id": "source-2026-09-09-005",
      "title": "arXiv preprint 2609.06063",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.06063",
      "canonical_url": "https://arxiv.org/abs/2609.06063",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-05T08:45:56.000-04:00",
      "accessed_at": "2026-09-09T08:30:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-09-005-001"
      ]
    },
    {
      "source_id": "source-2026-09-09-006",
      "title": "arXiv preprint 2609.06079",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.06079",
      "canonical_url": "https://arxiv.org/abs/2609.06079",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-05T09:13:07.000-04:00",
      "accessed_at": "2026-09-09T08:30:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-09-006-001"
      ]
    },
    {
      "source_id": "source-2026-09-09-007",
      "title": "arXiv preprint 2609.06128",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.06128",
      "canonical_url": "https://arxiv.org/abs/2609.06128",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-05T10:52:35.000-04:00",
      "accessed_at": "2026-09-09T08:30:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-09-007-001"
      ]
    },
    {
      "source_id": "source-2026-09-09-008",
      "title": "arXiv preprint 2609.06192",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.06192",
      "canonical_url": "https://arxiv.org/abs/2609.06192",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-05T13:29:15.000-04:00",
      "accessed_at": "2026-09-09T08:30:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-09-008-001"
      ]
    },
    {
      "source_id": "source-2026-09-09-009",
      "title": "arXiv preprint 2609.06831",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.06831",
      "canonical_url": "https://arxiv.org/abs/2609.06831",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-06T17:02:47.000-04:00",
      "accessed_at": "2026-09-09T08:30:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-09-009-001"
      ]
    },
    {
      "source_id": "source-2026-09-09-010",
      "title": "arXiv preprint 2609.06914",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.06914",
      "canonical_url": "https://arxiv.org/abs/2609.06914",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-06T21:30:34.000-04:00",
      "accessed_at": "2026-09-09T08:30:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-09-010-001"
      ]
    },
    {
      "source_id": "source-2026-09-09-011",
      "title": "arXiv preprint 2609.06960",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.06960",
      "canonical_url": "https://arxiv.org/abs/2609.06960",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-06T22:56:42.000-04:00",
      "accessed_at": "2026-09-09T08:30:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-09-011-001"
      ]
    },
    {
      "source_id": "source-2026-09-09-012",
      "title": "arXiv preprint 2609.06961",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.06961",
      "canonical_url": "https://arxiv.org/abs/2609.06961",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-06T22:59:56.000-04:00",
      "accessed_at": "2026-09-09T08:30:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-09-012-001"
      ]
    },
    {
      "source_id": "source-2026-09-09-013",
      "title": "arXiv preprint 2609.07032",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.07032",
      "canonical_url": "https://arxiv.org/abs/2609.07032",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-07T00:36:39.000-04:00",
      "accessed_at": "2026-09-09T08:30:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-09-013-001"
      ]
    },
    {
      "source_id": "source-2026-09-09-014",
      "title": "arXiv preprint 2609.07035",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.07035",
      "canonical_url": "https://arxiv.org/abs/2609.07035",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-07T00:44:19.000-04:00",
      "accessed_at": "2026-09-09T08:30:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-09-014-001"
      ]
    },
    {
      "source_id": "source-2026-09-09-015",
      "title": "arXiv preprint 2609.07057",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.07057",
      "canonical_url": "https://arxiv.org/abs/2609.07057",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-07T01:26:02.000-04:00",
      "accessed_at": "2026-09-09T08:30:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-09-026-001"
      ]
    },
    {
      "source_id": "source-2026-09-09-016",
      "title": "arXiv preprint 2609.07102",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.07102",
      "canonical_url": "https://arxiv.org/abs/2609.07102",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-07T02:44:01.000-04:00",
      "accessed_at": "2026-09-09T08:30:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-09-027-001"
      ]
    },
    {
      "source_id": "source-2026-09-09-017",
      "title": "arXiv preprint 2609.07075",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.07075",
      "canonical_url": "https://arxiv.org/abs/2609.07075",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-07T02:03:07.000-04:00",
      "accessed_at": "2026-09-09T08:30:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-09-015-001"
      ]
    },
    {
      "source_id": "source-2026-09-09-018",
      "title": "arXiv preprint 2609.07095",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.07095",
      "canonical_url": "https://arxiv.org/abs/2609.07095",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-07T02:33:04.000-04:00",
      "accessed_at": "2026-09-09T08:30:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-09-016-001"
      ]
    },
    {
      "source_id": "source-2026-09-09-019",
      "title": "arXiv preprint 2609.06820",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.06820",
      "canonical_url": "https://arxiv.org/abs/2609.06820",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-06T16:29:17.000-04:00",
      "accessed_at": "2026-09-09T08:30:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-09-017-001"
      ]
    },
    {
      "source_id": "source-2026-09-09-020",
      "title": "arXiv preprint 2609.07047",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.07047",
      "canonical_url": "https://arxiv.org/abs/2609.07047",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-07T01:06:20.000-04:00",
      "accessed_at": "2026-09-09T08:30:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-09-018-001"
      ]
    },
    {
      "source_id": "source-2026-09-09-021",
      "title": "arXiv preprint 2609.07176",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.07176",
      "canonical_url": "https://arxiv.org/abs/2609.07176",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-07T04:07:26.000-04:00",
      "accessed_at": "2026-09-09T08:30:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-09-019-001"
      ]
    }
  ],
  "corrections": [],
  "publisher": {
    "name": "The Machine Press",
    "url": "https://themachinepress.com",
    "description": "A daily newspaper for the age of artificial intelligence."
  }
}
