[MUSA][9/N] Add FA3 attention backend support through MATE (MUSA AI Tensor Engine) (#22051)
Co-authored-by: zhiguo.qin <zhiguo.qin@mthreads.com>
This commit is contained in:
@@ -114,8 +114,11 @@ srt_musa = [
|
||||
"sglang[runtime_common]",
|
||||
"torch",
|
||||
"torch_musa",
|
||||
"torchada>=0.1.45",
|
||||
"torchada>=0.1.48",
|
||||
"mthreads-ml-py",
|
||||
"mate",
|
||||
"mate-deep_gemm",
|
||||
"mate-flash-attention",
|
||||
"numpy<2.0",
|
||||
]
|
||||
|
||||
|
||||
Reference in New Issue
Block a user