chore: reduce layout threshold to 0.2, tune vLLM memory, and add test scripts and reports
This commit is contained in:
1 parent
bdb3a49742
commit
5c7c64e2f4
6 files changed
+2090
-9
No files matched your search
@@ -25,7 +25,7 @@ SubModules:
|
||||
model_name: PP-DocLayoutV3
|
||||
model_dir: null
|
||||
batch_size: 8
|
||||
threshold: 0.3
|
||||
threshold: 0.2
|
||||
layout_nms: True
|
||||
layout_unclip_ratio: [1.0, 1.0]
|
||||
layout_merge_bboxes_mode:
|
||||
|
||||
@@ -1,6 +1,6 @@
|
||||
# vLLM backend tuning for paddleocr genai_server
|
||||
# Docs: https://www.paddleocr.ai/latest/en/version3.x/pipeline_usage/PaddleOCR-VL.html#331-server-side-parameter-adjustment
|
||||
gpu-memory-utilization: 0.6
|
||||
gpu-memory-utilization: 0.35
|
||||
max-num-seqs: 4
|
||||
enforce-eager: true
|
||||
max-model-len: 2048
|
||||
|
||||
Reference in new issue
Block a user