mirror of
https://git.datalinker.icu/vllm-project/vllm.git
synced 2026-09-09 03:57:00 +08:00
[XPU][Triton]add xpu config in triton_reshape_and_cache_flash (#25643)
Signed-off-by: Kunshang Ji <kunshang.ji@intel.com> Signed-off-by: yewentao256 <zhyanwentao@126.com>
This commit is contained in:
parent
222411313d
commit
d7f6489f50
@ -137,7 +137,7 @@ def triton_reshape_and_cache_flash(
|
|||||||
|
|
||||||
# heuristics instead of autotuning
|
# heuristics instead of autotuning
|
||||||
TILE_SIZE = min(2048, triton.next_power_of_2(n))
|
TILE_SIZE = min(2048, triton.next_power_of_2(n))
|
||||||
if torch.version.hip:
|
if torch.version.hip or torch.version.xpu:
|
||||||
num_stages = 4
|
num_stages = 4
|
||||||
num_warps = 8
|
num_warps = 8
|
||||||
else: # cuda
|
else: # cuda
|
||||||
|
|||||||
Loading…
x
Reference in New Issue
Block a user