max_input_tokens 128000→120000 (GLM gateway unstable on very long inputs, 30s timeout at 128k)

Co-Authored-By: Claude <noreply@anthropic.com>
This commit is contained in:
sora 2026-09-11 09:22:55 +00:00
parent ebc51de274
commit b75fbc1a12

View File

@ -3,73 +3,52 @@ default:
top_p: 1.0
stream: true
max_tokens: 32768
aime24:
temperature: 1.0
repeats: 12
aime25:
temperature: 1.0
repeats: 12
aime26:
temperature: 1.0
repeats: 12
hmmt26:
temperature: 1.0
repeats: 12
imo_answerbench:
temperature: 1.0
gpqa_diamond:
temperature: 1.0
max_tokens: 8192
mmlu:
max_tokens: 8192
mmlu_pro:
max_tokens: 8192
cmmlu:
max_tokens: 8192
arc:
max_tokens: 8192
hellaswag:
max_tokens: 8192
winogrande:
max_tokens: 8192
simple_qa:
max_tokens: 8192
trivia_qa:
max_tokens: 8192
humaneval:
temperature: 1.0
live_code_bench:
temperature: 1.0
longbench_v2:
max_tokens: 8192
max_input_tokens: 128000
max_input_tokens: 120000
openai_mrcr:
max_tokens: 8192
max_input_tokens: 128000
max_input_tokens: 120000
bfcl_v3:
max_tokens: 4096
general_fc:
max_tokens: 4096
tau2_bench:
max_tokens: 16384