From d32e283947bed05378b1dec8f7322099e2c75aba Mon Sep 17 00:00:00 2001 From: amote-i <49533125+amote-i@users.noreply.github.com> Date: Fri, 8 May 2026 17:08:15 +0800 Subject: [PATCH] [NPU] [DOC] refresh npu supported model list (#24676) --- .../ascend-npus/ascend_npu_support_models.mdx | 110 +++++++++++++++--- 1 file changed, 94 insertions(+), 16 deletions(-) diff --git a/docs_new/docs/hardware-platforms/ascend-npus/ascend_npu_support_models.mdx b/docs_new/docs/hardware-platforms/ascend-npus/ascend_npu_support_models.mdx index d7ef03691..f3c05b12c 100644 --- a/docs_new/docs/hardware-platforms/ascend-npus/ascend_npu_support_models.mdx +++ b/docs_new/docs/hardware-platforms/ascend-npus/ascend_npu_support_models.mdx @@ -50,56 +50,104 @@ You are welcome to enable various models based on your business requirements. ✅ - Qwen/Qwen3.5-397B-A17B - Qwen + Eco-Tech/Qwen3.6-35B-A3B-w8a8 + Qwen3.6 + ✅ + ✅ + + + Eco-Tech/Qwen3.6-27B-w8a8 + Qwen3.6 + ✅ + ✅ + + + Eco-Tech/Qwen3.5-397B-A17B-w8a8-mtp + Qwen3.5 + ✅ + ✅ + + + Eco-Tech/Qwen3.5-122B-A10B-w8a8-mtp + Qwen3.5 + ✅ + ✅ + + + Eco-Tech/Qwen3.5-35B-A3B-w8a8-mtp + Qwen3.5 + ✅ + ✅ + + + Eco-Tech/Qwen3.5-27B-w8a8-mtp + Qwen3.5 + ✅ + ✅ + + + Qwen/Qwen3.5-9B + Qwen3.5 + ✅ + ✅ + + + Qwen/Qwen3.5-4B + Qwen3.5 + ✅ + ✅ + + + Qwen/Qwen3.5-0.8B + Qwen3.5 ✅ ✅ Qwen/Qwen3-30B-A3B-Instruct-2507 - Qwen + Qwen3 ✅ ✅ Qwen/Qwen3-32B - Qwen + Qwen3 ✅ ✅ Qwen/Qwen3-0.6B - Qwen + Qwen3 ✅ ✅ Qwen3-235B-A22B-W8A8 - Qwen + Qwen3 ✅ ✅ Qwen/Qwen3-Next-80B-A3B-Instruct - Qwen + Qwen3 ✅ ✅ Qwen3-Coder-480B-A35B-Instruct-w8a8-QuaRot - Qwen + Qwen3 ✅ ✅ Qwen/Qwen2.5-7B-Instruct - Qwen + Qwen2.5 ✅ ✅ QWQ-32B-W8A8 - Qwen + QWQ ✅ ✅ @@ -199,6 +247,18 @@ You are welcome to enable various models based on your business requirements. ✅ ✅ + + Eco-Tech/GLM-5.1-w4a8 + GLM-5.1 + ✅ + ✅ + + + Eco-Tech/GLM-5-w4a8 + GLM-5 + ✅ + ✅ + ZhipuAI/glm-4-9b-chat GLM-4 @@ -265,6 +325,18 @@ You are welcome to enable various models based on your business requirements. ✅ ✅ + + Eco-Tech/Kimi-K2.6-w4a8 + Kimi + ✅ + ✅ + + + Eco-Tech/Kimi-K2.5-w4a8 + Kimi + ✅ + ✅ + moonshotai/Kimi-K2-Thinking Kimi @@ -289,6 +361,12 @@ You are welcome to enable various models based on your business requirements. ✅ ✅ + + Eco-Tech/MiniMax-M2.5-w8a8-QuaRot + MiniMax-M2.5 + ✅ + ✅ + cyankiwi/MiniMax-M2-BF16 MiniMax-M2 @@ -354,37 +432,37 @@ You are welcome to enable various models based on your business requirements. Qwen/Qwen2.5-VL-3B-Instruct - Qwen-VL + Qwen2.5-VL ✅ ✅ Qwen/Qwen2.5-VL-72B-Instruct - Qwen-VL + Qwen2.5-VL ✅ ✅ Qwen/Qwen3-VL-30B-A3B-Instruct - Qwen-VL + Qwen3-VL ✅ ✅ Qwen/Qwen3-VL-8B-Instruct - Qwen-VL + Qwen3-VL ✅ ✅ Qwen/Qwen3-VL-4B-Instruct - Qwen-VL + Qwen3-VL ✅ ✅ Qwen/Qwen3-VL-235B-A22B-Instruct - Qwen-VL + Qwen3-VL ✅ ✅