{
  "$schema": "https://themachinepress.com/schemas/story-v1.schema.json",
  "schema_version": "1.0.0",
  "document_type": "machine_press_story",
  "story": {
    "story_id": "mp-2026-09-11-009",
    "source_story_id": "tmp-story-huro-robotized-human-video",
    "edition_id": "mp-2026-09-11-morning-0064",
    "edition_url": "https://themachinepress.com/edition/2026-09-11",
    "position": 9,
    "story_type": "dispatch",
    "section": "robotics",
    "editorial_classification": "editorial",
    "headline": "Six Hundred Thirty Thousand Human Videos Learned a Robot’s Shape",
    "slug": "six-hundred-thirty-thousand-human-videos-learned-a-robot-s-shape",
    "dek": "Robotizing observations and actions raised real-world completion from 51.5% to 80.3% as pretraining scale increased.",
    "summary": "Robotizing observations and actions raised real-world completion from 51.5% to 80.3% as pretraining scale increased.",
    "body_text": "HuRo converts heterogeneous human videos into robot-aligned observations and action trajectories, inferring missing intermediate signals across annotation levels. The resulting dataset contains roughly 630,000 episodes and 142 million processed frames from five video sources. Across four real manipulation tasks, scaling the robotized pretraining data increased overall completion from 51.5% to 80.3%; out-of-distribution completion under spatial and visual shifts rose from 34.9% to 72.2%. Ablations attribute some robustness to visual robotization and favor end-to-end retargeted actions over visual-only transfer. The evidence covers the authors’ four tasks, not arbitrary robots or internet video.",
    "why_it_matters": "Robotizing observations and actions raised real-world completion from 51.5% to 80.3% as pretraining scale increased.",
    "limitations": [
      "Ablations attribute some robustness to visual robotization and favor end-to-end retargeted actions over visual-only transfer.",
      "The evidence covers the authors’ four tasks, not arbitrary robots or internet video."
    ],
    "importance": 9,
    "canonical_url": "https://themachinepress.com/story/mp-2026-09-11-009/six-hundred-thirty-thousand-human-videos-learned-a-robot-s-shape",
    "json_url": "https://themachinepress.com/story/mp-2026-09-11-009.json",
    "first_published_at": "2026-09-11T09:00:00.000-04:00",
    "modified_at": "2026-09-11T09:00:00.000-04:00",
    "content_status": "new",
    "is_carryover": false,
    "carryover_reason": null,
    "key_claims": [
      {
        "claim_id": "claim-mp-2026-09-11-009-001",
        "text": "Robotizing observations and actions raised real-world completion from 51.5% to 80.3% as pretraining scale increased.",
        "source_ids": [
          "source-2026-09-11-009"
        ],
        "qualification": "Ablations attribute some robustness to visual robotization and favor end-to-end retargeted actions over visual-only transfer."
      }
    ],
    "source_ids": [
      "source-2026-09-11-009"
    ],
    "tags": [
      "robot learning",
      "human video",
      "VLA"
    ],
    "image_url": null,
    "corrections": []
  },
  "sources": [
    {
      "source_id": "source-2026-09-11-009",
      "title": "arXiv preprint 2609.10706",
      "publisher": "arXiv",
      "url": "https://arxiv.org/abs/2609.10706",
      "canonical_url": "https://arxiv.org/abs/2609.10706",
      "source_type": "primary_research",
      "is_primary_source": true,
      "published_at": "2026-09-09T14:02:05.000-04:00",
      "accessed_at": "2026-09-11T08:25:00.000-04:00",
      "supports_claim_ids": [
        "claim-mp-2026-09-11-009-001"
      ]
    }
  ],
  "corrections": [],
  "publisher": {
    "name": "The Machine Press",
    "url": "https://themachinepress.com",
    "description": "A daily newspaper for the age of artificial intelligence."
  },
  "cite_this_report": {
    "title": "Six Hundred Thirty Thousand Human Videos Learned a Robot’s Shape",
    "publisher": "The Machine Press",
    "published_at": "2026-09-11T09:00:00.000-04:00",
    "canonical_url": "https://themachinepress.com/story/mp-2026-09-11-009/six-hundred-thirty-thousand-human-videos-learned-a-robot-s-shape"
  }
}
