mirror of
https://git.datalinker.icu/vllm-project/vllm.git
synced 2025-12-10 01:35:01 +08:00
8 lines
588 B
Plaintext
8 lines
588 B
Plaintext
compressed-tensors, nm-testing/Mixtral-8x7B-Instruct-v0.1-W4A16-quantized, main
|
|
compressed-tensors, nm-testing/Mixtral-8x7B-Instruct-v0.1-W4A16-channel-quantized, main
|
|
compressed-tensors, nm-testing/Mixtral-8x7B-Instruct-v0.1-W8A16-quantized, main
|
|
compressed-tensors, nm-testing/test-w4a16-mixtral-actorder-group, main
|
|
gptq_marlin, TheBloke/Mixtral-8x7B-v0.1-GPTQ, main
|
|
gptq_marlin, TheBloke/Mixtral-8x7B-v0.1-GPTQ, gptq-8bit-128g-actorder_True
|
|
awq_marlin, casperhansen/deepseek-coder-v2-instruct-awq, main
|
|
compressed-tensors, RedHatAI/Llama-4-Scout-17B-16E-Instruct-quantized.w4a16, main |