File size: 1,428 Bytes
b53331c
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
2
3
4
5
6
7
8
9
10
11
12
13
14
15
16
17
18
19
20
21
22
23
24
25
26
27
28
29
30
31
32
33
34
{
  "status": "PASS",
  "stage": "B14 REAP-Code-S1 selected-only accumulation validation",
  "decision_eligible": true,
  "release_eligible": false,
  "model_path": "/workspace/models/glm-4.7-flash",
  "model_revision": "7dd20894a642a0aa287e9827cb1a1f7f91386b67",
  "model_class": "Glm4MoeLiteForCausalLM",
  "model_dtype": "torch.bfloat16",
  "experts_implementation": "eager",
  "manifest_path": "/workspace/kguard-reap/data/glm-calibration-reap-code-s1-8192.jsonl",
  "manifest_sha256": "3f2c46de59e42213b9d4c614cd9fdaca3d53b5ee141cf13c13680aee71cc326b",
  "samples": 8192,
  "shards": 128,
  "shard_size": 64,
  "total_tokens": 4366438,
  "collection_mode": "batch1-all-token",
  "observed_moe_layer_count": 46,
  "experts": 64,
  "top_k": 4,
  "scaling": 1.8,
  "observer_logit_max_error_this_run": 0.0,
  "load_seconds_this_run": 372.36645814100484,
  "peak_cuda_allocated_gib_this_run": 59.07405138015747,
  "collection_seconds_all_shards": 7233.351361623034,
  "tokens_per_second_all_shards": 603.653518501312,
  "tensor_output_path": "/workspace/kguard-reap/logs/glm-reap-code-s1-combined-8192.pt",
  "tensor_output_sha256": "ec5959ed72a50e04c23809190d22632fb9b4dcb77d4c6f590c7f581b78ece148",
  "detail_output_path": "/workspace/kguard-reap/logs/glm-reap-code-s1-combined-8192.json",
  "detail_output_sha256": "83e33f906ce9617f8052d0f33fdc602a58cfff882b99e5065248628e24cb7f38",
  "failure_count": 0,
  "failures": []
}