{"ok":true,"entity":{"id":"vision-language-action-model","name":"Vision-Language-Action Model","entityType":"concept","officialName":"Vision-Language-Action Model","canonicalName":"Vision-Language-Action Model","displayName":"Vision-Language-Action Model","category":"Vision-Language-Actionモデル（概念）","shortDescription":"視覚入力と自然言語による指示を受け取り、ロボットの低レベルな動作コマンドを直接出力するマルチモーダル基盤モデルの一種。RT-2、OpenVLA、Physical Intelligenceのpi0などが代表例とされる。","primaryCluster":"physical-ai-robotics","parentEntity":null,"verificationStatus":"draft","website":null,"updatedAt":"2026-09-01T12:27:03.812Z","secondaryClusters":[],"alias":["VLA model","VLA","Vision Language Action model"],"searchKeywords":["VLAモデル","ロボット基盤モデル","pi0","OpenVLA","マルチモーダルAI"]},"referenceIndex":["P-01-001","P-02-001","P-04-001","P-05-001","P-03-001"]}