[CI Sprint] Quantization CI Cleanup (#24130)

Signed-off-by: Alex Yun <alexyun04@gmail.com>
This commit is contained in:
Alex
2025-11-18 08:21:48 -06:00
committed by GitHub
parent 184b12fdc6
commit f6aa122698
10 changed files with 32 additions and 26 deletions

View File

@@ -23,8 +23,8 @@ from vllm.model_executor.layers.quantization import (
get_quantization_config,
register_quantization_config,
)
from vllm.model_executor.layers.quantization.base_config import ( # noqa: E501
QuantizationConfig,
from vllm.model_executor.layers.quantization.base_config import (
QuantizationConfig, # noqa: E501
)
@@ -142,5 +142,5 @@ def test_custom_quant(vllm_runner, model, monkeypatch):
llm.apply_model(check_model)
output = llm.generate_greedy("Hello my name is", max_tokens=20)
output = llm.generate_greedy("Hello my name is", max_tokens=1)
assert output