Reserve more memory for DeepSeekOCR model and adjust server start timeout for DeepGEMM to reduce flakiness (#15277)

This commit is contained in:
Kangyan-Zhou
2025-12-17 13:13:05 -08:00
committed by GitHub
parent 5290cef97c
commit 011d8d8970
3 changed files with 20 additions and 7 deletions
+4
View File
@@ -162,6 +162,10 @@ class TestQwen2AudioServer(AudioOpenAITestMixin):
class TestDeepseekOCRServer(TestOpenAIMLLMServerBase):
model = "deepseek-ai/DeepSeek-OCR"
trust_remote_code = False
extra_args = [
"--mem-fraction-static=0.70",
"--cuda-graph-max-bs=4",
]
def verify_single_image_response_for_ocr(self, response):
"""Verify DeepSeek-OCR grounding output with coordinates"""