patdev commited on
Commit
e4c5faa
·
verified ·
1 Parent(s): 61c9ef7

journal de demarrage

Browse files
Files changed (1) hide show
  1. etat/xtn8ts2meqsqb4.log +10 -1
etat/xtn8ts2meqsqb4.log CHANGED
@@ -1,4 +1,4 @@
1
- === bootstrap v86-journal-runtime-continu | 00:10:05 UTC ===
2
  [VL] 23:36:01 bootstrap v86-journal-runtime-continu
3
  [VL] 23:36:01 pilote 595.91.07, CUDA runtime 13.2
4
  [VL] 23:36:01 pilote 595.91.07 : compat CUDA non necessaire
@@ -295,6 +295,13 @@
295
  (APIServer pid=5703) INFO 09-06 00:08:23 [loggers.py:310] Engine 000: Avg prompt throughput: 0.0 tokens/s, Avg generation throughput: 9.5 tokens/s, Running: 0 reqs, Waiting: 0 reqs, GPU KV cache usage: 0.0%, Prefix cache hit rate: 19.0%
296
  (APIServer pid=5703) INFO 09-06 00:08:23 [metrics.py:120] SpecDecoding metrics: Mean acceptance length: 1.00, Accepted throughput: 0.00 tokens/s, Drafted throughput: 28.50 tokens/s, Accepted: 0 tokens, Drafted: 285 tokens, Per-position acceptance rate: 0.000, 0.000, 0.000, Avg Draft acceptance rate: 0.0%
297
  (APIServer pid=5703) INFO 09-06 00:08:33 [loggers.py:310] Engine 000: Avg prompt throughput: 0.0 tokens/s, Avg generation throughput: 0.0 tokens/s, Running: 1 reqs, Waiting: 0 reqs, GPU KV cache usage: 10.3%, Prefix cache hit rate: 20.3%
 
 
 
 
 
 
 
298
 
299
  === proxy.log (fin) ===
300
  INFO: Started server process [6819]
@@ -317,3 +324,5 @@ INFO: 100.64.1.5:49334 - "POST /v1/messages?beta=true HTTP/1.1" 200 OK
317
  INFO: 100.64.1.2:43518 - "POST /v1/messages?beta=true HTTP/1.1" 200 OK
318
  INFO: 100.64.1.1:54690 - "POST /v1/messages?beta=true HTTP/1.1" 200 OK
319
  INFO: 100.64.1.3:38406 - "POST /v1/messages?beta=true HTTP/1.1" 200 OK
 
 
 
1
+ === bootstrap v86-journal-runtime-continu | 00:12:08 UTC ===
2
  [VL] 23:36:01 bootstrap v86-journal-runtime-continu
3
  [VL] 23:36:01 pilote 595.91.07, CUDA runtime 13.2
4
  [VL] 23:36:01 pilote 595.91.07 : compat CUDA non necessaire
 
295
  (APIServer pid=5703) INFO 09-06 00:08:23 [loggers.py:310] Engine 000: Avg prompt throughput: 0.0 tokens/s, Avg generation throughput: 9.5 tokens/s, Running: 0 reqs, Waiting: 0 reqs, GPU KV cache usage: 0.0%, Prefix cache hit rate: 19.0%
296
  (APIServer pid=5703) INFO 09-06 00:08:23 [metrics.py:120] SpecDecoding metrics: Mean acceptance length: 1.00, Accepted throughput: 0.00 tokens/s, Drafted throughput: 28.50 tokens/s, Accepted: 0 tokens, Drafted: 285 tokens, Per-position acceptance rate: 0.000, 0.000, 0.000, Avg Draft acceptance rate: 0.0%
297
  (APIServer pid=5703) INFO 09-06 00:08:33 [loggers.py:310] Engine 000: Avg prompt throughput: 0.0 tokens/s, Avg generation throughput: 0.0 tokens/s, Running: 1 reqs, Waiting: 0 reqs, GPU KV cache usage: 10.3%, Prefix cache hit rate: 20.3%
298
+ (APIServer pid=5703) INFO 09-06 00:10:23 [loggers.py:310] Engine 000: Avg prompt throughput: 24239.5 tokens/s, Avg generation throughput: 31.4 tokens/s, Running: 1 reqs, Waiting: 0 reqs, GPU KV cache usage: 24.1%, Prefix cache hit rate: 20.3%
299
+ (APIServer pid=5703) INFO 09-06 00:10:23 [metrics.py:120] SpecDecoding metrics: Mean acceptance length: 1.00, Accepted throughput: 0.00 tokens/s, Drafted throughput: 7.82 tokens/s, Accepted: 0 tokens, Drafted: 939 tokens, Per-position acceptance rate: 0.000, 0.000, 0.000, Avg Draft acceptance rate: 0.0%
300
+ (APIServer pid=5703) INFO: 172.71.134.75:0 - "POST /v1/messages/count_tokens?beta=true HTTP/1.1" 200 OK
301
+ (APIServer pid=5703) INFO: 172.71.134.75:0 - "POST /v1/messages?beta=true HTTP/1.1" 200 OK
302
+ (APIServer pid=5703) INFO 09-06 00:10:33 [loggers.py:310] Engine 000: Avg prompt throughput: 0.0 tokens/s, Avg generation throughput: 13.5 tokens/s, Running: 0 reqs, Waiting: 0 reqs, GPU KV cache usage: 0.0%, Prefix cache hit rate: 20.3%
303
+ (APIServer pid=5703) INFO 09-06 00:10:33 [metrics.py:120] SpecDecoding metrics: Mean acceptance length: 1.00, Accepted throughput: 0.00 tokens/s, Drafted throughput: 40.50 tokens/s, Accepted: 0 tokens, Drafted: 405 tokens, Per-position acceptance rate: 0.000, 0.000, 0.000, Avg Draft acceptance rate: 0.0%
304
+ (APIServer pid=5703) INFO 09-06 00:10:43 [loggers.py:310] Engine 000: Avg prompt throughput: 0.0 tokens/s, Avg generation throughput: 0.0 tokens/s, Running: 1 reqs, Waiting: 0 reqs, GPU KV cache usage: 10.3%, Prefix cache hit rate: 21.3%
305
 
306
  === proxy.log (fin) ===
307
  INFO: Started server process [6819]
 
324
  INFO: 100.64.1.2:43518 - "POST /v1/messages?beta=true HTTP/1.1" 200 OK
325
  INFO: 100.64.1.1:54690 - "POST /v1/messages?beta=true HTTP/1.1" 200 OK
326
  INFO: 100.64.1.3:38406 - "POST /v1/messages?beta=true HTTP/1.1" 200 OK
327
+ INFO: 100.64.1.3:38406 - "POST /v1/messages/count_tokens?beta=true HTTP/1.1" 200 OK
328
+ INFO: 100.64.1.3:38406 - "POST /v1/messages?beta=true HTTP/1.1" 200 OK