{"ok":true,"entity":{"id":"sglang","name":"SGLang","entityType":"product","officialName":"SGLang","canonicalName":"SGLang","displayName":"SGLang","category":"推論エンジン（LLM配信）","shortDescription":"構造化されたLLMプログラムを高速に実行するための推論エンジン。RadixAttentionによるKVキャッシュ再利用等の最適化により、単一GPUから大規模分散クラスタまで高スループットな配信を実現する。","primaryCluster":"ai-open-model","parentEntity":null,"verificationStatus":"draft","website":"https://github.com/sgl-project/sglang","updatedAt":"2026-07-20T02:14:59.356Z","secondaryClusters":["ai-infrastructure"],"alias":["SGLang Runtime"],"searchKeywords":["LLM推論エンジン","RadixAttention","構造化生成","高速推論"]},"referenceIndex":["P-01-001","P-02-001","P-04-001","P-03-001","P-06-001"]}