[Kernel] Relocate vendored fla and mamba kernel trees to sglang.kernels (RFC #29630, Phase 2.5, 7/7) (#30795)
Co-authored-by: Claude Fable 5 <noreply@anthropic.com>
This commit is contained in:
co-authored by
Claude Fable 5
parent
23f2b77d82
commit
4aadf94146
@@ -7,7 +7,7 @@ import torch.utils.benchmark as benchmark
|
||||
from flashinfer import BatchDecodeWithPagedKVCacheWrapper
|
||||
|
||||
from sglang.kernels.ops.attention.decode_attention import decode_attention_fwd
|
||||
from sglang.srt.layers.attention.flashinfer_backend import should_use_tensor_core
|
||||
from sglang.kernels.ops.attention.flashinfer_backend import should_use_tensor_core
|
||||
|
||||
|
||||
def benchmark_forward(
|
||||
|
||||
Reference in New Issue
Block a user