chore: reduce layout threshold to 0.2, tune vLLM memory, and add test scripts and reports

This commit is contained in:
Rafhan Mazaya Fathurrahman committed 2026-07-02 15:30:19 +07:00
1 parent bdb3a49742
commit 5c7c64e2f4
6 files changed
+2090 -9

No files matched your search

+1 -1
View File
@@ -1,6 +1,6 @@
# vLLM backend tuning for paddleocr genai_server
# Docs: https://www.paddleocr.ai/latest/en/version3.x/pipeline_usage/PaddleOCR-VL.html#331-server-side-parameter-adjustment
gpu-memory-utilization: 0.6
gpu-memory-utilization: 0.35
max-num-seqs: 4
enforce-eager: true
max-model-len: 2048