{"ok":true,"entity":{"slug":"sglang","entityType":"product","name":"SGLang","officialName":"SGLang","canonicalName":"SGLang","displayName":"SGLang","category":"推論エンジン（LLM配信）","shortDescription":"構造化されたLLMプログラムを高速に実行するための推論エンジン。RadixAttentionによるKVキャッシュ再利用等の最適化により、単一GPUから大規模分散クラスタまで高スループットな配信を実現する。","alias":["SGLang Runtime"],"searchKeywords":["LLM推論エンジン","RadixAttention","構造化生成","高速推論"],"website":"https://github.com/sgl-project/sglang","parentEntity":null,"primaryCluster":"ai-open-model","secondaryClusters":["ai-infrastructure"],"id":"sglang","verificationStatus":"draft","updatedAt":"2026-07-20T02:14:59.356Z"},"referenceIndex":["P-01-001","P-02-001","P-04-001"]}