{"ok":true,"entity":{"id":"reinforcement-learning","name":"強化学習（Reinforcement Learning）","entityType":"concept","canonicalName":"Reinforcement Learning","displayName":"強化学習（Reinforcement Learning）","category":"AI基礎概念","shortDescription":"エージェントが環境と相互作用し、報酬を最大化するよう試行錯誤で方策を学ぶ機械学習の枠組み。","primaryCluster":"ai-company","parentEntity":null,"verificationStatus":"draft","website":null,"updatedAt":"2026-07-10T00:42:57.294Z","secondaryClusters":[],"alias":["RL","強化学習"],"searchKeywords":["reinforcement learning","RL","報酬"]},"referenceIndex":["P-01-001","P-02-001","P-04-001","P-01-002","P-04-002","P-05-001"]}