{
  "$schema": "https://themachinepress.com/schemas/edition-v1.schema.json",
  "schema_version": "1.0.0",
  "document_type": "machine_press_edition",
  "edition_id": "mp-2026-08-23-morning-0045",
  "source_edition_id": "mp-2026-08-23-morning-0045",
  "edition_number": 45,
  "edition_label": "Morning edition",
  "status": "published",
  "revision": 1,
  "published_at": "2026-08-23T09:00:00.000-04:00",
  "modified_at": "2026-08-23T09:00:00.000-04:00",
  "timezone": "America/Detroit",
  "canonical_url": "https://themachinepress.com/edition/2026-08-23",
  "html_url": "https://themachinepress.com/edition/2026-08-23",
  "json_url": "https://themachinepress.com/edition/2026-08-23.json",
  "markdown_url": "https://themachinepress.com/edition/2026-08-23.md",
  "lead_story_id": "mp-2026-08-23-001",
  "lead_summary": "An executed-replay audit found that three common step-level credit signals for tool-using agents identified causal contribution no better than chance.",
  "coverage_window": {
    "start": "2026-08-20T01:41:23.000-04:00",
    "end": "2026-08-23T08:27:00.000-04:00"
  },
  "story_count": 27,
  "stories": [
    {
      "story_id": "mp-2026-08-23-001",
      "source_story_id": "tmp-lead-agent-credit-executed-replay",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 1,
      "story_type": "lead",
      "section": "benchmarks-evals",
      "editorial_classification": "editorial",
      "headline": "The Judge Couldn’t Find the Step That Changed the Outcome",
      "slug": "the-judge-couldn-t-find-the-step-that-changed-the-outcome",
      "dek": "An executed-replay audit found that three common step-level credit signals for tool-using agents identified causal contribution no better than chance.",
      "summary": "An executed-replay audit found that three common step-level credit signals for tool-using agents identified causal contribution no better than chance.",
      "body_text": "The researchers resampled a policy’s own alternatives at each decision point in ALFWorld and rolled the trajectory forward, creating an executed counterfactual measure of what actually changed the outcome. Against that causal reference, LLM-judge scores, outcome-conditioned log-probability ratios and the policy’s confidence all performed at chance; the authors also report that measurable contribution was sparse and that the available counterfactuals depended on the policy. Their seven-arm training experiment found no arm that reliably beat the untrained policy, while differing sample counts explained apparent training signatures. The result is a preprint finding in one single-agent environment, but its warning is broader: a fluent-looking training signal can measure exposure or correctness without identifying the step that caused success.",
      "why_it_matters": "An executed-replay audit found that three common step-level credit signals for tool-using agents identified causal contribution no better than chance.",
      "limitations": [],
      "importance": 10,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-001/the-judge-couldn-t-find-the-step-that-changed-the-outcome",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-001.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-08-23-001-001",
          "text": "An executed-replay audit found that three common step-level credit signals for tool-using agents identified causal contribution no better than chance.",
          "source_ids": [
            "source-2026-08-23-001"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-08-23-001"
      ],
      "tags": [
        "AI agents",
        "credit assignment",
        "executed replay",
        "evaluation"
      ],
      "image_url": "https://themachinepress.com/issues/2026-08-23/lead-agent-credit-replay.png",
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-002",
      "source_story_id": "tmp-feature-bedroom-radar-comparison",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 2,
      "story_type": "secondary",
      "section": "research",
      "editorial_classification": "editorial",
      "headline": "The Ceiling Saw Sleep Without a Camera",
      "slug": "the-ceiling-saw-sleep-without-a-camera",
      "dek": "FMCW radar, ultra-wideband radar and Wi-Fi sensing were recorded together across twenty people and six room layouts.",
      "summary": "FMCW radar, ultra-wideband radar and Wi-Fi sensing were recorded together across twenty people and six room layouts.",
      "body_text": "A controlled study mounted three contact-free radio systems above the same bedroom scenes and evaluated them with the same convolutional network. IR-UWB produced the strongest cross-subject result on the ten-class activity task, while FMCW generalized best to unseen room layouts; all three technologies exceeded 92 percent macro F1 on the study’s four-class sleep-monitoring task in unseen environments. The authors attribute the trade-off to differences in range resolution, antenna diversity, Doppler resolution and retained spatial information. The experiment involved twenty participants and should not be read as clinical validation, but it gives designers a rare like-for-like comparison instead of forcing them to compare results gathered with different rooms, hardware and methods.",
      "why_it_matters": "FMCW radar, ultra-wideband radar and Wi-Fi sensing were recorded together across twenty people and six room layouts.",
      "limitations": [
        "The experiment involved twenty participants and should not be read as clinical validation, but it gives designers a rare like-for-like comparison instead of forcing them to compare results gathered with different rooms, hardware and methods."
      ],
      "importance": 9,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-002/the-ceiling-saw-sleep-without-a-camera",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-002.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-08-23-002-001",
          "text": "FMCW radar, ultra-wideband radar and Wi-Fi sensing were recorded together across twenty people and six room layouts.",
          "source_ids": [
            "source-2026-08-23-002"
          ],
          "qualification": "The experiment involved twenty participants and should not be read as clinical validation, but it gives designers a rare like-for-like comparison instead of forcing them to compare results gathered with different rooms, hardware and methods."
        }
      ],
      "source_ids": [
        "source-2026-08-23-002"
      ],
      "tags": [
        "radar",
        "Wi-Fi sensing",
        "sleep monitoring",
        "health technology"
      ],
      "image_url": "https://themachinepress.com/issues/2026-08-23/feature-bedroom-radar.png",
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-003",
      "source_story_id": "tmp-story-agent-memory-evolving-state",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 3,
      "story_type": "dispatch",
      "section": "developer-tools",
      "editorial_classification": "editorial",
      "headline": "The Memory Kept the Old Fact After the World Changed",
      "slug": "the-memory-kept-the-old-fact-after-the-world-changed",
      "dek": "StateMemBench separates current-state answers from superseded facts across 234 multi-session scenarios.",
      "summary": "StateMemBench separates current-state answers from superseded facts across 234 multi-session scenarios.",
      "body_text": "The benchmark tests whether an agent updates its working world when facts, constraints and decisions change, rather than merely retrieving something that was once true. The authors report that a state-first method improved current-state accuracy over same-backbone and existing-memory baselines, and that a single-call wrapper added 32 to 67 points across six backends; absolute benchmark performance remained limited, keeping the result squarely in research territory.",
      "why_it_matters": "StateMemBench separates current-state answers from superseded facts across 234 multi-session scenarios.",
      "limitations": [],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-003/the-memory-kept-the-old-fact-after-the-world-changed",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-003.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-08-23-003-001",
          "text": "StateMemBench separates current-state answers from superseded facts across 234 multi-session scenarios.",
          "source_ids": [
            "source-2026-08-23-003"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-08-23-003"
      ],
      "tags": [
        "agent memory",
        "state tracking",
        "benchmarks"
      ],
      "image_url": "https://themachinepress.com/issues/2026-08-23/agent-memory-server-file-image.webp",
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-004",
      "source_story_id": "tmp-story-fleetsieve-llm-profiling",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 4,
      "story_type": "dispatch",
      "section": "chips-infrastructure",
      "editorial_classification": "editorial",
      "headline": "The Fleet Stopped Profiling When the Decision Stopped Moving",
      "slug": "the-fleet-stopped-profiling-when-the-decision-stopped-moving",
      "dek": "FleetSieve measures only the LLM-serving configurations likely to change a resource-coupled allocation.",
      "summary": "FleetSieve measures only the LLM-serving configurations likely to change a resource-coupled allocation.",
      "body_text": "On a fixed H100 grid for a 31-billion-parameter open model, the method matched the oracle aggregate choice with 6.9 percent fewer GPU-seconds than uniform random profiling in the fixed comparison. Its joint capacity and tail-latency model also avoided a configuration whose 46.4-second p99 violated a 30-second service objective, although the paper reports that FleetSieve did not use the fewest GPU-seconds for every workload.",
      "why_it_matters": "FleetSieve measures only the LLM-serving configurations likely to change a resource-coupled allocation.",
      "limitations": [
        "Its joint capacity and tail-latency model also avoided a configuration whose 46.4-second p99 violated a 30-second service objective, although the paper reports that FleetSieve did not use the fewest GPU-seconds for every workload."
      ],
      "importance": 7,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-004/the-fleet-stopped-profiling-when-the-decision-stopped-moving",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-004.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-08-23-004-001",
          "text": "FleetSieve measures only the LLM-serving configurations likely to change a resource-coupled allocation.",
          "source_ids": [
            "source-2026-08-23-004"
          ],
          "qualification": "Its joint capacity and tail-latency model also avoided a configuration whose 46.4-second p99 violated a 30-second service objective, although the paper reports that FleetSieve did not use the fewest GPU-seconds for every workload."
        }
      ],
      "source_ids": [
        "source-2026-08-23-004"
      ],
      "tags": [
        "LLM serving",
        "GPU fleets",
        "latency"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-005",
      "source_story_id": "tmp-story-recache-agent-kv-blocks",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 5,
      "story_type": "dispatch",
      "section": "chips-infrastructure",
      "editorial_classification": "editorial",
      "headline": "The Tool Schema Became a Reusable Cache Block",
      "slug": "the-tool-schema-became-a-reusable-cache-block",
      "dek": "ReCache separates recurring tool and skill descriptions so their key-value states survive changes in order and combination.",
      "summary": "ReCache separates recurring tool and skill descriptions so their key-value states survive changes in order and combination.",
      "body_text": "Resource-local attention and positions make cached schema blocks composition-invariant, while route selection and pruning reduce what remains visible at inference. On a benchmark assembled from seven public tool-and-skill datasets, the preprint reports nearly unchanged invocation F1, a 3.655-times time-to-first-token speedup for resource-wise attention and a 92.43 percent reduction in allocated KV-tensor memory for the complete framework.",
      "why_it_matters": "ReCache separates recurring tool and skill descriptions so their key-value states survive changes in order and combination.",
      "limitations": [],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-005/the-tool-schema-became-a-reusable-cache-block",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-005.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-08-23-005-001",
          "text": "ReCache separates recurring tool and skill descriptions so their key-value states survive changes in order and combination.",
          "source_ids": [
            "source-2026-08-23-005"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-08-23-005"
      ],
      "tags": [
        "KV cache",
        "tool use",
        "inference"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-006",
      "source_story_id": "tmp-story-asymmetric-llm-compression-harms",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 6,
      "story_type": "dispatch",
      "section": "safety-security",
      "editorial_classification": "editorial",
      "headline": "Compression Hid the Knowledge It Lost",
      "slug": "compression-hid-the-knowledge-it-lost",
      "dek": "Aggregate accuracy and bias scores concealed subgroup shifts and confident errors across eleven compression methods.",
      "summary": "Aggregate accuracy and bias scores concealed subgroup shifts and confident errors across eleven compression methods.",
      "body_text": "Researchers evaluated three language models across eleven compression methods and found that compressed systems disproportionately lost head knowledge relative to tail knowledge while often remaining confident about newly incorrect answers. Stable aggregate bias scores also masked opposing movements across demographic subgroups, arguing for granular deployment audits rather than a single perplexity, accuracy or bias number.",
      "why_it_matters": "Aggregate accuracy and bias scores concealed subgroup shifts and confident errors across eleven compression methods.",
      "limitations": [],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-006/compression-hid-the-knowledge-it-lost",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-006.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-08-23-006-001",
          "text": "Aggregate accuracy and bias scores concealed subgroup shifts and confident errors across eleven compression methods.",
          "source_ids": [
            "source-2026-08-23-006"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-08-23-006"
      ],
      "tags": [
        "model compression",
        "bias",
        "knowledge retention"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-007",
      "source_story_id": "tmp-story-thinkingbox-stateful-workflows",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 7,
      "story_type": "dispatch",
      "section": "benchmarks-evals",
      "editorial_classification": "editorial",
      "headline": "Passing Once Collapsed to 25 Percent Across Twenty Runs",
      "slug": "passing-once-collapsed-to-25-percent-across-twenty-runs",
      "dek": "Thinkingbox evaluates terminal backend state, policy compliance and collateral effects across 507 business workflows.",
      "summary": "Thinkingbox evaluates terminal backend state, policy compliance and collateral effects across 507 business workflows.",
      "body_text": "The strongest tested model reached 65.36 percent pass-at-one but only 25.25 percent success across twenty attempts on the new benchmark. Many failed runs terminated cleanly after valid state-changing actions, supporting the authors’ argument that a plausible response or tool call is not evidence that the right persistent state transition occurred without extra effects.",
      "why_it_matters": "Thinkingbox evaluates terminal backend state, policy compliance and collateral effects across 507 business workflows.",
      "limitations": [
        "The strongest tested model reached 65.36 percent pass-at-one but only 25.25 percent success across twenty attempts on the new benchmark.",
        "Many failed runs terminated cleanly after valid state-changing actions, supporting the authors’ argument that a plausible response or tool call is not evidence that the right persistent state transition occurred without extra effects."
      ],
      "importance": 9,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-007/passing-once-collapsed-to-25-percent-across-twenty-runs",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-007.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-08-23-007-001",
          "text": "Thinkingbox evaluates terminal backend state, policy compliance and collateral effects across 507 business workflows.",
          "source_ids": [
            "source-2026-08-23-007"
          ],
          "qualification": "The strongest tested model reached 65.36 percent pass-at-one but only 25.25 percent success across twenty attempts on the new benchmark."
        }
      ],
      "source_ids": [
        "source-2026-08-23-007"
      ],
      "tags": [
        "agent reliability",
        "business workflows",
        "MCP"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-008",
      "source_story_id": "tmp-story-goag-object-agnostic-grasping",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 8,
      "story_type": "dispatch",
      "section": "robotics",
      "editorial_classification": "editorial",
      "headline": "The Gripper Learned Its Own Surface Before Seeing the Object",
      "slug": "the-gripper-learned-its-own-surface-before-seeing-the-object",
      "dek": "GOAG introduces object features only at inference time and samples contacts from a learned representation of the hand.",
      "summary": "GOAG introduces object features only at inference time and samples contacts from a learned representation of the hand.",
      "body_text": "The generative planner starts from the geometric fact that gripper and object surfaces coincide at valid contacts, then models the contact distribution for a specific gripper without object-specific training data. The authors report an 86.93 percent average success rate on MultiDex objects plus simulated and real-world tests across multiple grippers; the claim remains tied to the reported protocols, not universal dexterity.",
      "why_it_matters": "GOAG introduces object features only at inference time and samples contacts from a learned representation of the hand.",
      "limitations": [
        "The authors report an 86.93 percent average success rate on MultiDex objects plus simulated and real-world tests across multiple grippers; the claim remains tied to the reported protocols, not universal dexterity."
      ],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-008/the-gripper-learned-its-own-surface-before-seeing-the-object",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-008.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-08-23-008-001",
          "text": "GOAG introduces object features only at inference time and samples contacts from a learned representation of the hand.",
          "source_ids": [
            "source-2026-08-23-008"
          ],
          "qualification": "The authors report an 86.93 percent average success rate on MultiDex objects plus simulated and real-world tests across multiple grippers; the claim remains tied to the reported protocols, not universal dexterity."
        }
      ],
      "source_ids": [
        "source-2026-08-23-008"
      ],
      "tags": [
        "robotics",
        "grasp planning",
        "generative models"
      ],
      "image_url": "https://themachinepress.com/issues/2026-08-23/goag-robot-file-image.webp",
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-009",
      "source_story_id": "tmp-story-swe-bench-science",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 9,
      "story_type": "dispatch",
      "section": "benchmarks-evals",
      "editorial_classification": "editorial",
      "headline": "The Best Coding Agent Solved Fewer Than Half",
      "slug": "the-best-coding-agent-solved-fewer-than-half",
      "dek": "SWE-bench Science spans 119 tasks, 98 repositories and twenty fields where software is part of the instrument.",
      "summary": "SWE-bench Science spans 119 tasks, 98 repositories and twenty fields where software is part of the instrument.",
      "body_text": "The benchmark separates issue-driven, expert-exploratory and engineering-integration work and reports a best pass-at-one below 50 percent. Its error analysis finds failures in scientific abstraction, exploration, repair coverage and generalization; a paired ablation also showed that well-grounded scientific guidance can help while poorly aligned guidance can anchor the repair in the wrong direction.",
      "why_it_matters": "SWE-bench Science spans 119 tasks, 98 repositories and twenty fields where software is part of the instrument.",
      "limitations": [],
      "importance": 9,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-009/the-best-coding-agent-solved-fewer-than-half",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-009.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-08-23-009-001",
          "text": "SWE-bench Science spans 119 tasks, 98 repositories and twenty fields where software is part of the instrument.",
          "source_ids": [
            "source-2026-08-23-009"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-08-23-009"
      ],
      "tags": [
        "coding agents",
        "scientific software",
        "benchmarks"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-010",
      "source_story_id": "tmp-story-llm-judge-panel-routing",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 10,
      "story_type": "dispatch",
      "section": "benchmarks-evals",
      "editorial_classification": "editorial",
      "headline": "The Panel Stopped Calling Copies of the Same Judge",
      "slug": "the-panel-stopped-calling-copies-of-the-same-judge",
      "dek": "A role-conditioned allocation method drops redundant judges, routes specialists by slice and stops when validation gain saturates.",
      "summary": "A role-conditioned allocation method drops redundant judges, routes specialists by slice and stops when validation gain saturates.",
      "body_text": "The method uses a small labeled audit set, declared slices and call costs to distinguish copies, global complements and conditional specialists. Across reasoning, code, safety, preference, reward, summarization and math audits, the output is an auditable call plan rather than a claim that one fixed panel wins everywhere.",
      "why_it_matters": "A role-conditioned allocation method drops redundant judges, routes specialists by slice and stops when validation gain saturates.",
      "limitations": [],
      "importance": 7,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-010/the-panel-stopped-calling-copies-of-the-same-judge",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-010.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-08-23-010-001",
          "text": "A role-conditioned allocation method drops redundant judges, routes specialists by slice and stops when validation gain saturates.",
          "source_ids": [
            "source-2026-08-23-010"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-08-23-010"
      ],
      "tags": [
        "LLM judges",
        "routing",
        "evaluation cost"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-011",
      "source_story_id": "tmp-story-industrial-abnormal-situation-reasoning",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 11,
      "story_type": "dispatch",
      "section": "safety",
      "editorial_classification": "editorial",
      "headline": "The Reasoning Model Kept All 39 Plant Scenarios Inside Bounds",
      "slug": "the-reasoning-model-kept-all-39-plant-scenarios-inside-bounds",
      "dek": "A programmatically bounded action interface put general-purpose reasoning models against a plant-wide benchmark.",
      "summary": "A programmatically bounded action interface put general-purpose reasoning models against a plant-wide benchmark.",
      "body_text": "The authors report that the leading model maintained hard constraints across all 39 abnormal situations and operating-point changes, while basic regulatory control failed in fifteen. It also matched the benchmark’s expert-engineered advanced control and diagnosed the stated root cause in fifteen safety-critical cases; these are benchmark results without a human in the loop, not authorization to deploy an LLM on a real plant.",
      "why_it_matters": "A programmatically bounded action interface put general-purpose reasoning models against a plant-wide benchmark.",
      "limitations": [
        "It also matched the benchmark’s expert-engineered advanced control and diagnosed the stated root cause in fifteen safety-critical cases; these are benchmark results without a human in the loop, not authorization to deploy an LLM on a real plant."
      ],
      "importance": 9,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-011/the-reasoning-model-kept-all-39-plant-scenarios-inside-bounds",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-011.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-08-23-011-001",
          "text": "A programmatically bounded action interface put general-purpose reasoning models against a plant-wide benchmark.",
          "source_ids": [
            "source-2026-08-23-011"
          ],
          "qualification": "It also matched the benchmark’s expert-engineered advanced control and diagnosed the stated root cause in fifteen safety-critical cases; these are benchmark results without a human in the loop, not authorization to deploy an LLM on a real plant."
        }
      ],
      "source_ids": [
        "source-2026-08-23-011"
      ],
      "tags": [
        "industrial control",
        "safety",
        "reasoning models"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-012",
      "source_story_id": "tmp-story-inadvertent-context-leakage",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 12,
      "story_type": "dispatch",
      "section": "safety-security",
      "editorial_classification": "editorial",
      "headline": "A Benign Reply Carried Four Digits of the Secret",
      "slug": "a-benign-reply-carried-four-digits-of-the-secret",
      "dek": "Researchers reconstructed in-context secrets from ordinary outputs even when models refused direct extraction.",
      "summary": "Researchers reconstructed in-context secrets from ordinary outputs even when models refused direct extraction.",
      "body_text": "Across eight proprietary models in controlled experiments, the authors report near-perfect recovery of two-digit secrets and 82 percent exact recovery for four digits from benign responses. They also trained attacks that infer predicates about user memories and extract longer identifiers in a production-style agent, framing context sensitivity itself as a covert leakage channel that capability may amplify.",
      "why_it_matters": "Researchers reconstructed in-context secrets from ordinary outputs even when models refused direct extraction.",
      "limitations": [],
      "importance": 9,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-012/a-benign-reply-carried-four-digits-of-the-secret",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-012.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-08-23-012-001",
          "text": "Researchers reconstructed in-context secrets from ordinary outputs even when models refused direct extraction.",
          "source_ids": [
            "source-2026-08-23-012"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-08-23-012"
      ],
      "tags": [
        "privacy",
        "context windows",
        "agent security"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-013",
      "source_story_id": "tmp-story-malicious-skill-benchmark",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 13,
      "story_type": "dispatch",
      "section": "safety-security",
      "editorial_classification": "editorial",
      "headline": "The Detector Recognized the Source More Than the Threat",
      "slug": "the-detector-recognized-the-source-more-than-the-threat",
      "dek": "MaliciousSkillBench consolidates 9,740 agent-skill packages and tests whether detection transfers beyond familiar feeds.",
      "summary": "MaliciousSkillBench consolidates 9,740 agent-skill packages and tests whether detection transfers beyond familiar feeds.",
      "body_text": "The benchmark reduces 8,414 raw malicious records to 7,539 normalized identities and evaluates both learned detectors and off-the-shelf scanners. Random-split macro F1 reached as high as 0.932, but source-disjoint performance fell to 0.653–0.665; the strongest text model retained high malicious recall while falsely flagging 62.4 percent of benign skills from held-out sources.",
      "why_it_matters": "MaliciousSkillBench consolidates 9,740 agent-skill packages and tests whether detection transfers beyond familiar feeds.",
      "limitations": [],
      "importance": 9,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-013/the-detector-recognized-the-source-more-than-the-threat",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-013.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-08-23-013-001",
          "text": "MaliciousSkillBench consolidates 9,740 agent-skill packages and tests whether detection transfers beyond familiar feeds.",
          "source_ids": [
            "source-2026-08-23-013"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-08-23-013"
      ],
      "tags": [
        "agent skills",
        "malware detection",
        "benchmarks"
      ],
      "image_url": "https://themachinepress.com/issues/2026-08-23/malicious-skills-file-image.webp",
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-014",
      "source_story_id": "tmp-story-shieldfs-confidential-storage",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 14,
      "story_type": "dispatch",
      "section": "infrastructure",
      "editorial_classification": "editorial",
      "headline": "The Trusted Enclave Still Needed a Fresh Disk",
      "slug": "the-trusted-enclave-still-needed-a-fresh-disk",
      "dek": "ShieldFS extends ZFS so a hostile storage stack cannot silently roll back, replay or fork persistent state.",
      "summary": "ShieldFS extends ZFS so a hostile storage stack cannot silently roll back, replay or fork persistent state.",
      "body_text": "The design keeps succinct commitments inside trusted execution environments and a lightweight registry, while authenticating the write-ahead log and storage pool with hash chains and an embedded Merkle tree. Reads verify freshness and integrity without application changes; the paper reports performance comparable to other evaluated filesystems, but the security claim remains conditioned on its confidential-computing threat model and trusted registry.",
      "why_it_matters": "ShieldFS extends ZFS so a hostile storage stack cannot silently roll back, replay or fork persistent state.",
      "limitations": [],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-014/the-trusted-enclave-still-needed-a-fresh-disk",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-014.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-08-23-014-001",
          "text": "ShieldFS extends ZFS so a hostile storage stack cannot silently roll back, replay or fork persistent state.",
          "source_ids": [
            "source-2026-08-23-014"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-08-23-014"
      ],
      "tags": [
        "confidential computing",
        "filesystems",
        "storage security"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-026",
      "source_story_id": "tmp-story-policyguide-agent-workflows",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 15,
      "story_type": "dispatch",
      "section": "business-enterprise",
      "editorial_classification": "editorial",
      "headline": "The Guardrail Became a Workflow",
      "slug": "the-guardrail-became-a-workflow",
      "dek": "PolicyGuide compiles policy into a graph and returns step-specific remediation at user-turn boundaries.",
      "summary": "PolicyGuide compiles policy into a graph and returns step-specific remediation at user-turn boundaries.",
      "body_text": "Across airline, retail and telecom tasks with one agent-verifier pairing, the preprint reports mean four-run pass rate rising from 0.42 to 0.62, with the largest improvement in telecom. The same workflow representation transferred to two other agent families, while complementary tests found lower observed attack success and stronger procedural compliance than the compared safeguards.",
      "why_it_matters": "PolicyGuide compiles policy into a graph and returns step-specific remediation at user-turn boundaries.",
      "limitations": [],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-026/the-guardrail-became-a-workflow",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-026.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-08-23-026-001",
          "text": "PolicyGuide compiles policy into a graph and returns step-specific remediation at user-turn boundaries.",
          "source_ids": [
            "source-2026-08-23-015"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-08-23-015"
      ],
      "tags": [
        "policy compliance",
        "customer service",
        "agent workflows"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-027",
      "source_story_id": "tmp-story-evidence-arbitration-text-numbers",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 16,
      "story_type": "dispatch",
      "section": "benchmarks-evals",
      "editorial_classification": "editorial",
      "headline": "The Model Followed Recency More Than Reliability",
      "slug": "the-model-followed-recency-more-than-reliability",
      "dek": "A controlled benchmark made textual summaries, numerical series and external forecasts disagree on purpose.",
      "summary": "A controlled benchmark made textual summaries, numerical series and external forecasts disagree on purpose.",
      "body_text": "Because the synthetic risk trajectories identify which evidence source matches the ground truth, the study can vary modality, recency, stated reliability and provenance independently. Open-weight instruction models showed systematic text-versus-number preferences and followed recent evidence more consistently than reliability labels, sometimes over-weighting an external forecast against direct context. The finding isolates a heuristic failure mode for tool-augmented decisions rather than measuring a deployed domain.",
      "why_it_matters": "A controlled benchmark made textual summaries, numerical series and external forecasts disagree on purpose.",
      "limitations": [],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-027/the-model-followed-recency-more-than-reliability",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-027.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-08-23-027-001",
          "text": "A controlled benchmark made textual summaries, numerical series and external forecasts disagree on purpose.",
          "source_ids": [
            "source-2026-08-23-016"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-08-23-016"
      ],
      "tags": [
        "evidence arbitration",
        "tool use",
        "reliability"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-015",
      "source_story_id": "tmp-sidebar-speech-synthesizer-deception",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 17,
      "story_type": "ticker",
      "section": "safety-security",
      "editorial_classification": "editorial",
      "headline": "One Altered Sentence Fooled Listeners 77 Percent of the Time",
      "slug": "one-altered-sentence-fooled-listeners-77-percent-of-the-time",
      "dek": "In a study of 82 IT professionals, partial audio spoofs were harder to localize than fully synthetic speech.",
      "summary": "In a study of 82 IT professionals, partial audio spoofs were harder to localize than fully synthetic speech.",
      "body_text": "Strict accuracy for locating the altered sentence fell to 9 percent, and listeners marked the synthetic segment as genuine 77 percent of the time. Humans and six pretrained detectors failed in different ways, supporting procedural verification and provenance rather than listening alone.",
      "why_it_matters": "In a study of 82 IT professionals, partial audio spoofs were harder to localize than fully synthetic speech.",
      "limitations": [],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-015/one-altered-sentence-fooled-listeners-77-percent-of-the-time",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-015.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-08-23-015-001",
          "text": "In a study of 82 IT professionals, partial audio spoofs were harder to localize than fully synthetic speech.",
          "source_ids": [
            "source-2026-08-23-017"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-08-23-017"
      ],
      "tags": [
        "audio deepfakes",
        "speech synthesis",
        "provenance"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-016",
      "source_story_id": "tmp-sidebar-green-compression-break-even",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 18,
      "story_type": "ticker",
      "section": "research",
      "editorial_classification": "editorial",
      "headline": "Compression Needed a Carbon Break-Even Ledger",
      "slug": "compression-needed-a-carbon-break-even-ledger",
      "dek": "Two internship projects compare the footprint of ML training and inference with storage saved by lossless compression.",
      "summary": "Two internship projects compare the footprint of ML training and inference with storage saved by lossless compression.",
      "body_text": "The note frames environmental benefit as a break-even calculation instead of assuming that fewer stored bytes are automatically greener. It is a concise project report, not a universal lifecycle estimate.",
      "why_it_matters": "Two internship projects compare the footprint of ML training and inference with storage saved by lossless compression.",
      "limitations": [
        "It is a concise project report, not a universal lifecycle estimate."
      ],
      "importance": 6,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-016/compression-needed-a-carbon-break-even-ledger",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-016.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-08-23-016-001",
          "text": "Two internship projects compare the footprint of ML training and inference with storage saved by lossless compression.",
          "source_ids": [
            "source-2026-08-23-018"
          ],
          "qualification": "It is a concise project report, not a universal lifecycle estimate."
        }
      ],
      "source_ids": [
        "source-2026-08-23-018"
      ],
      "tags": [
        "compression",
        "carbon accounting",
        "storage"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-017",
      "source_story_id": "tmp-sidebar-evidence-gated-robot-planning",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 19,
      "story_type": "ticker",
      "section": "robotics",
      "editorial_classification": "editorial",
      "headline": "The Robot Looked Before Planning With a Missing Object",
      "slug": "the-robot-looked-before-planning-with-a-missing-object",
      "dek": "EAFG acquires visual evidence, proceeds, explores again or halts before long-horizon manipulation.",
      "summary": "EAFG acquires visual evidence, proceeds, explores again or halts before long-horizon manipulation.",
      "body_text": "In ambiguous cooking tasks, the framework found task-relevant objects before planning and reduced repeated attempts when a required object was absent. The result targets a specific partial-observability failure in vision-language task-and-motion planning.",
      "why_it_matters": "EAFG acquires visual evidence, proceeds, explores again or halts before long-horizon manipulation.",
      "limitations": [],
      "importance": 7,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-017/the-robot-looked-before-planning-with-a-missing-object",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-017.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-08-23-017-001",
          "text": "EAFG acquires visual evidence, proceeds, explores again or halts before long-horizon manipulation.",
          "source_ids": [
            "source-2026-08-23-019"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-08-23-019"
      ],
      "tags": [
        "robot planning",
        "visual evidence",
        "partial observability"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-018",
      "source_story_id": "tmp-sidebar-verified-compiler-performance",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 20,
      "story_type": "ticker",
      "section": "developer-tools",
      "editorial_classification": "editorial",
      "headline": "The Optimizer Got a Compile-Time Bound",
      "slug": "the-optimizer-got-a-compile-time-bound",
      "dek": "A Rocq proof covers semantic preservation, monotone improvement, convergence time and output performance for inlining.",
      "summary": "A Rocq proof covers semantic preservation, monotone improvement, convergence time and output performance for inlining.",
      "body_text": "The proof-of-concept treats compiler performance and unpredictable search time as properties worth verifying alongside semantics. It applies a cache-cost model to inline expansion rather than claiming a verified optimizer for every pass.",
      "why_it_matters": "A Rocq proof covers semantic preservation, monotone improvement, convergence time and output performance for inlining.",
      "limitations": [],
      "importance": 7,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-018/the-optimizer-got-a-compile-time-bound",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-018.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-08-23-018-001",
          "text": "A Rocq proof covers semantic preservation, monotone improvement, convergence time and output performance for inlining.",
          "source_ids": [
            "source-2026-08-23-020"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-08-23-020"
      ],
      "tags": [
        "compilers",
        "formal verification",
        "inlining"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-019",
      "source_story_id": "tmp-sidebar-phantom-self-improvement",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 21,
      "story_type": "ticker",
      "section": "benchmarks-evals",
      "editorial_classification": "editorial",
      "headline": "The Frozen Model Invented Improvements Too",
      "slug": "the-frozen-model-invented-improvements-too",
      "dek": "A measured null exposed seven ways transition-level self-improvement ledgers can manufacture gains.",
      "summary": "A measured null exposed seven ways transition-level self-improvement ledgers can manufacture gains.",
      "body_text": "A frozen control pushed through the same evaluation pipeline showed apparent capability changes, including artifacts tied to greedy decoding and batching. A per-problem exact test with false-discovery-rate control found no changes on held-out null replicates, making matched baseline measurement the paper’s central recommendation.",
      "why_it_matters": "A measured null exposed seven ways transition-level self-improvement ledgers can manufacture gains.",
      "limitations": [],
      "importance": 8,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-019/the-frozen-model-invented-improvements-too",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-019.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [
        {
          "claim_id": "claim-mp-2026-08-23-019-001",
          "text": "A measured null exposed seven ways transition-level self-improvement ledgers can manufacture gains.",
          "source_ids": [
            "source-2026-08-23-021"
          ],
          "qualification": null
        }
      ],
      "source_ids": [
        "source-2026-08-23-021"
      ],
      "tags": [
        "self-improvement",
        "evaluation",
        "measurement null"
      ],
      "image_url": null,
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-020",
      "source_story_id": "invention-tulip-creative-computer",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 22,
      "story_type": "invention_desk",
      "section": "invention-desk",
      "editorial_classification": "editorial",
      "headline": "Tulip Creative Computer",
      "slug": "tulip-creative-computer",
      "dek": "A self-contained touchscreen computer boots into MicroPython and combines programmable graphics, MIDI connections, and an open-source synthesizer in focused, buildable hardware.",
      "summary": "A self-contained touchscreen computer boots into MicroPython and combines programmable graphics, MIDI connections, and an open-source synthesizer in focused, buildable hardware.",
      "body_text": "A self-contained touchscreen computer boots into MicroPython and combines programmable graphics, MIDI connections, and an open-source synthesizer in focused, buildable hardware.",
      "why_it_matters": "An independent builder is turning an improbable idea into a working project.",
      "limitations": [
        "A Desk Pick is an editorial selection, not a product endorsement."
      ],
      "importance": null,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-020/tulip-creative-computer",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-020.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [],
      "source_ids": [],
      "tags": [
        "independent-builders",
        "desk-pick",
        "released"
      ],
      "image_url": "https://themachinepress.com/issues/2026-08-23/invention-desk/tulip-creative-computer.webp",
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-021",
      "source_story_id": "invention-open-press-project",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 23,
      "story_type": "invention_desk",
      "section": "invention-desk",
      "editorial_classification": "editorial",
      "headline": "Open Press Project",
      "slug": "open-press-project",
      "dek": "Shrinks an intaglio press into 3D-printed tabletop hardware, with free fabrication plans for makers and finished presses for artists without a printer.",
      "summary": "Shrinks an intaglio press into 3D-printed tabletop hardware, with free fabrication plans for makers and finished presses for artists without a printer.",
      "body_text": "Shrinks an intaglio press into 3D-printed tabletop hardware, with free fabrication plans for makers and finished presses for artists without a printer.",
      "why_it_matters": "An independent builder is turning an improbable idea into a working project.",
      "limitations": [
        "A Desk Pick is an editorial selection, not a product endorsement."
      ],
      "importance": null,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-021/open-press-project",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-021.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [],
      "source_ids": [],
      "tags": [
        "independent-builders",
        "desk-pick",
        "released"
      ],
      "image_url": "https://themachinepress.com/issues/2026-08-23/invention-desk/open-press-project.webp",
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-022",
      "source_story_id": "invention-openflexure-microscope",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 24,
      "story_type": "invention_desk",
      "section": "invention-desk",
      "editorial_classification": "editorial",
      "headline": "OpenFlexure Microscope",
      "slug": "openflexure-microscope",
      "dek": "Prints most of a precise microscope body as one flexure mechanism, pairing interchangeable optics with optional motorized sample positioning and open control software.",
      "summary": "Prints most of a precise microscope body as one flexure mechanism, pairing interchangeable optics with optional motorized sample positioning and open control software.",
      "body_text": "Prints most of a precise microscope body as one flexure mechanism, pairing interchangeable optics with optional motorized sample positioning and open control software.",
      "why_it_matters": "An independent builder is turning an improbable idea into a working project.",
      "limitations": [
        "A Desk Pick is an editorial selection, not a product endorsement."
      ],
      "importance": null,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-022/openflexure-microscope",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-022.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [],
      "source_ids": [],
      "tags": [
        "independent-builders",
        "desk-pick",
        "released"
      ],
      "image_url": "https://themachinepress.com/issues/2026-08-23/invention-desk/openflexure-microscope.webp",
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-023",
      "source_story_id": "invention-satnogs",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 25,
      "story_type": "invention_desk",
      "section": "invention-desk",
      "editorial_classification": "editorial",
      "headline": "SatNOGS",
      "slug": "satnogs",
      "dek": "Links volunteer-built radio ground stations to shared scheduling and observation services, turning backyard antennas into a public satellite-observation network.",
      "summary": "Links volunteer-built radio ground stations to shared scheduling and observation services, turning backyard antennas into a public satellite-observation network.",
      "body_text": "Links volunteer-built radio ground stations to shared scheduling and observation services, turning backyard antennas into a public satellite-observation network.",
      "why_it_matters": "An independent builder is turning an improbable idea into a working project.",
      "limitations": [
        "A Desk Pick is an editorial selection, not a product endorsement."
      ],
      "importance": null,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-023/satnogs",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-023.json",
      "first_published_at": "2026-08-23T09:00:00.000-04:00",
      "modified_at": "2026-08-23T09:00:00.000-04:00",
      "content_status": "new",
      "is_carryover": false,
      "carryover_reason": null,
      "key_claims": [],
      "source_ids": [],
      "tags": [
        "independent-builders",
        "desk-pick",
        "released"
      ],
      "image_url": "https://themachinepress.com/issues/2026-08-23/invention-desk/satnogs.webp",
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-024",
      "source_story_id": "invention-sponsored-house-example",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 26,
      "story_type": "invention_desk",
      "section": "invention-desk",
      "editorial_classification": "house_example",
      "headline": "The First Paid Slot",
      "slug": "the-first-paid-slot",
      "dek": "A transparent preview of paid placement with one verified link and no claim of endorsement.",
      "summary": "A transparent preview of paid placement with one verified link and no claim of endorsement.",
      "body_text": "A transparent preview of paid placement with one verified link and no claim of endorsement.\n\nHouse example - no advertiser paid. Payment will buy placement, never endorsement.",
      "why_it_matters": "This placement explains how builders can appear in The Invention Desk without purchasing editorial endorsement.",
      "limitations": [
        "House example - no advertiser paid. Payment will buy placement, never endorsement."
      ],
      "importance": null,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-024/the-first-paid-slot",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-024.json",
      "first_published_at": "2026-07-10T09:00:00.000-04:00",
      "modified_at": "2026-07-10T09:00:00.000-04:00",
      "content_status": "carried_over",
      "is_carryover": true,
      "carryover_reason": "invention_desk_seven_day_placement",
      "key_claims": [],
      "source_ids": [],
      "tags": [
        "independent-builders",
        "sponsored-project",
        "open"
      ],
      "image_url": "https://themachinepress.com/issues/2026-08-23/invention-desk/sponsored-project.webp",
      "corrections": []
    },
    {
      "story_id": "mp-2026-08-23-025",
      "source_story_id": "invention-placement-cta",
      "edition_id": "mp-2026-08-23-morning-0045",
      "edition_url": "https://themachinepress.com/edition/2026-08-23",
      "position": 27,
      "story_type": "invention_desk",
      "section": "invention-desk",
      "editorial_classification": "house_example",
      "headline": "Put Your Project on the Desk",
      "slug": "put-your-project-on-the-desk",
      "dek": "One manually reviewed placement stays active for seven days and remains separate from Desk Picks.",
      "summary": "One manually reviewed placement stays active for seven days and remains separate from Desk Picks.",
      "body_text": "One manually reviewed placement stays active for seven days and remains separate from Desk Picks.\n\nManual intake only. Payment buys placement, never endorsement, and every submission is reviewed.",
      "why_it_matters": "This placement explains how builders can appear in The Invention Desk without purchasing editorial endorsement.",
      "limitations": [
        "Manual intake only. Payment buys placement, never endorsement, and every submission is reviewed."
      ],
      "importance": null,
      "canonical_url": "https://themachinepress.com/story/mp-2026-08-23-025/put-your-project-on-the-desk",
      "json_url": "https://themachinepress.com/story/mp-2026-08-23-025.json",
      "first_published_at": "2026-07-10T09:00:00.000-04:00",
      "modified_at": "2026-07-10T09:00:00.000-04:00",
      "content_status": "carried_over",
      "is_carryover": true,
      "carryover_reason": "invention_desk_seven_day_placement",
      "key_claims": [],
      "source_ids": [],
      "tags": [
        "independent-builders",
        "cta",
        "open"
      ],
      "image_url": "https://themachinepress.com/issues/2026-08-23/invention-desk/put-your-project-on-the-desk.webp",
      "corrections": []
    }
  ],
  "sources": [
    {
      "source_id": "source-2026-08-23-001",
      "title": "arXiv preprint 2608.19760",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2608.19760",
      "canonical_url": "https://arxiv.org/abs/2608.19760",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-08-20T04:04:00.000-04:00",
      "accessed_at": "2026-08-23T08:27:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-08-23-001-001"
      ]
    },
    {
      "source_id": "source-2026-08-23-002",
      "title": "arXiv preprint 2608.20322",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2608.20322",
      "canonical_url": "https://arxiv.org/abs/2608.20322",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-08-20T13:58:22.000-04:00",
      "accessed_at": "2026-08-23T08:27:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-08-23-002-001"
      ]
    },
    {
      "source_id": "source-2026-08-23-003",
      "title": "arXiv preprint 2608.19652",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2608.19652",
      "canonical_url": "https://arxiv.org/abs/2608.19652",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-08-20T01:41:23.000-04:00",
      "accessed_at": "2026-08-23T08:27:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-08-23-003-001"
      ]
    },
    {
      "source_id": "source-2026-08-23-004",
      "title": "arXiv preprint 2608.19659",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2608.19659",
      "canonical_url": "https://arxiv.org/abs/2608.19659",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-08-20T01:54:29.000-04:00",
      "accessed_at": "2026-08-23T08:27:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-08-23-004-001"
      ]
    },
    {
      "source_id": "source-2026-08-23-005",
      "title": "arXiv preprint 2608.19662",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2608.19662",
      "canonical_url": "https://arxiv.org/abs/2608.19662",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-08-20T01:57:24.000-04:00",
      "accessed_at": "2026-08-23T08:27:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-08-23-005-001"
      ]
    },
    {
      "source_id": "source-2026-08-23-006",
      "title": "arXiv preprint 2608.19670",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2608.19670",
      "canonical_url": "https://arxiv.org/abs/2608.19670",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-08-20T02:06:14.000-04:00",
      "accessed_at": "2026-08-23T08:27:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-08-23-006-001"
      ]
    },
    {
      "source_id": "source-2026-08-23-007",
      "title": "arXiv preprint 2608.19741",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2608.19741",
      "canonical_url": "https://arxiv.org/abs/2608.19741",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-08-20T03:37:57.000-04:00",
      "accessed_at": "2026-08-23T08:27:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-08-23-007-001"
      ]
    },
    {
      "source_id": "source-2026-08-23-008",
      "title": "arXiv preprint 2608.19759",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2608.19759",
      "canonical_url": "https://arxiv.org/abs/2608.19759",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-08-20T04:03:39.000-04:00",
      "accessed_at": "2026-08-23T08:27:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-08-23-008-001"
      ]
    },
    {
      "source_id": "source-2026-08-23-009",
      "title": "arXiv preprint 2608.19799",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2608.19799",
      "canonical_url": "https://arxiv.org/abs/2608.19799",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-08-20T04:53:15.000-04:00",
      "accessed_at": "2026-08-23T08:27:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-08-23-009-001"
      ]
    },
    {
      "source_id": "source-2026-08-23-010",
      "title": "arXiv preprint 2608.19802",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2608.19802",
      "canonical_url": "https://arxiv.org/abs/2608.19802",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-08-20T04:58:00.000-04:00",
      "accessed_at": "2026-08-23T08:27:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-08-23-010-001"
      ]
    },
    {
      "source_id": "source-2026-08-23-011",
      "title": "arXiv preprint 2608.19819",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2608.19819",
      "canonical_url": "https://arxiv.org/abs/2608.19819",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-08-20T05:13:37.000-04:00",
      "accessed_at": "2026-08-23T08:27:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-08-23-011-001"
      ]
    },
    {
      "source_id": "source-2026-08-23-012",
      "title": "arXiv preprint 2608.19857",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2608.19857",
      "canonical_url": "https://arxiv.org/abs/2608.19857",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-08-20T06:05:29.000-04:00",
      "accessed_at": "2026-08-23T08:27:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-08-23-012-001"
      ]
    },
    {
      "source_id": "source-2026-08-23-013",
      "title": "arXiv preprint 2608.19901",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2608.19901",
      "canonical_url": "https://arxiv.org/abs/2608.19901",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-08-20T07:13:00.000-04:00",
      "accessed_at": "2026-08-23T08:27:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-08-23-013-001"
      ]
    },
    {
      "source_id": "source-2026-08-23-014",
      "title": "arXiv preprint 2608.19924",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2608.19924",
      "canonical_url": "https://arxiv.org/abs/2608.19924",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-08-20T07:41:15.000-04:00",
      "accessed_at": "2026-08-23T08:27:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-08-23-014-001"
      ]
    },
    {
      "source_id": "source-2026-08-23-015",
      "title": "arXiv preprint 2608.19861",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2608.19861",
      "canonical_url": "https://arxiv.org/abs/2608.19861",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-08-20T06:13:19.000-04:00",
      "accessed_at": "2026-08-23T08:27:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-08-23-026-001"
      ]
    },
    {
      "source_id": "source-2026-08-23-016",
      "title": "arXiv preprint 2608.20116",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2608.20116",
      "canonical_url": "https://arxiv.org/abs/2608.20116",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-08-20T10:48:30.000-04:00",
      "accessed_at": "2026-08-23T08:27:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-08-23-027-001"
      ]
    },
    {
      "source_id": "source-2026-08-23-017",
      "title": "arXiv preprint 2608.19959",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2608.19959",
      "canonical_url": "https://arxiv.org/abs/2608.19959",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-08-20T08:27:37.000-04:00",
      "accessed_at": "2026-08-23T08:27:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-08-23-015-001"
      ]
    },
    {
      "source_id": "source-2026-08-23-018",
      "title": "arXiv preprint 2608.19994",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2608.19994",
      "canonical_url": "https://arxiv.org/abs/2608.19994",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-08-20T09:08:51.000-04:00",
      "accessed_at": "2026-08-23T08:27:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-08-23-016-001"
      ]
    },
    {
      "source_id": "source-2026-08-23-019",
      "title": "arXiv preprint 2608.20084",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2608.20084",
      "canonical_url": "https://arxiv.org/abs/2608.20084",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-08-20T10:17:52.000-04:00",
      "accessed_at": "2026-08-23T08:27:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-08-23-017-001"
      ]
    },
    {
      "source_id": "source-2026-08-23-020",
      "title": "arXiv preprint 2608.20137",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2608.20137",
      "canonical_url": "https://arxiv.org/abs/2608.20137",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-08-20T11:02:11.000-04:00",
      "accessed_at": "2026-08-23T08:27:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-08-23-018-001"
      ]
    },
    {
      "source_id": "source-2026-08-23-021",
      "title": "arXiv preprint 2608.20290",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2608.20290",
      "canonical_url": "https://arxiv.org/abs/2608.20290",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-08-20T13:30:14.000-04:00",
      "accessed_at": "2026-08-23T08:27:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-08-23-019-001"
      ]
    }
  ],
  "corrections": [],
  "publisher": {
    "name": "The Machine Press",
    "url": "https://themachinepress.com",
    "description": "A daily newspaper for the age of artificial intelligence."
  }
}
