Fix torch version check for SM100 mxfp4 (#22535)

Signed-off-by: Zifei Tong <zifeitong@gmail.com> Signed-off-by: mgoin <mgoin64@gmail.com> Co-authored-by: mgoin <mgoin64@gmail.com>

Fix torch version check for SM100 mxfp4 (#22535)
Signed-off-by: Zifei Tong <zifeitong@gmail.com> Signed-off-by: mgoin <mgoin64@gmail.com> Co-authored-by: mgoin <mgoin64@gmail.com>
6534d2fc · zifeitong · GitHub · 422f22e0 · 6534d2fc
Unverified Commit 6534d2fc authored Aug 12, 2025 by zifeitong Committed by GitHub Aug 12, 2025
Hide whitespace changes
Inline Side-by-side

Showing with 8 additions and 6 deletions

vllm/model_executor/layers/fused_moe/layer.py vllm/model_executor/layers/fused_moe/layer.py +8 -6

No files found.
--- a/vllm/model_executor/layers/fused_moe/layer.py
+++ b/vllm/model_executor/layers/fused_moe/layer.py
@@ -741,12 +741,14 @@ class FusedMoE(torch.nn.Module):
        # we padding globally so EP buffer allocation works
        if quant_config and quant_config.get_name() == "mxfp4":
-            if not is_torch_equal_or_newer("2.8.0"):
+            if not current_platform.is_device_capability(100):
-                raise RuntimeError("Mxfp4 on hopper requires torch >= 2.8.0")
+                if not is_torch_equal_or_newer("2.8.0"):
-            if current_platform.is_device_capability(
+                    raise RuntimeError(
-                    90) and not has_triton_kernels():
+                        "Mxfp4 on non-blackwell requires torch >= 2.8.0")
-                raise NotImplementedError(
+                if not has_triton_kernels():
-                    "Triton kernels must be installed for mxfp4 on hopper")
+                    raise NotImplementedError(
+                        "triton_kernels must be installed for "
+                        "mxfp4 on non-blackwell")
            if (current_platform.is_rocm()
                    or envs.VLLM_USE_FLASHINFER_MOE_MXFP4_MXFP8
                    or envs.VLLM_USE_FLASHINFER_MOE_MXFP4_BF16):