| { |
| "case_name": "poly", |
| "peft_version": "0.20.0", |
| "git_commit": "a5e689189fd2dd79f67ad2adee7f2e9e6328c82f", |
| "torch_version": "2.13.0+cu130", |
| "transformers_version": "5.15.0.dev0", |
| "base_model_id": "peft-internal-testing/tiny-random-T5ForConditionalGeneration-calibrated", |
| "model_cls": "AutoModelForSeq2SeqLM", |
| "inputs": { |
| "input_ids": [ |
| [ |
| 1, |
| 2, |
| 3 |
| ], |
| [ |
| 6, |
| 5, |
| 4 |
| ] |
| ], |
| "attention_mask": [ |
| [ |
| 1, |
| 1, |
| 1 |
| ], |
| [ |
| 1, |
| 1, |
| 1 |
| ] |
| ], |
| "decoder_input_ids": [ |
| [ |
| 0, |
| 1, |
| 2 |
| ], |
| [ |
| 2, |
| 0, |
| 1 |
| ] |
| ], |
| "task_ids": [ |
| 0, |
| 1 |
| ] |
| }, |
| "atol": 1e-06, |
| "rtol": 1e-06, |
| "state_dict_keys": [ |
| "base_model.model.decoder.block.0.layer.0.SelfAttention.q.poly_lora_A", |
| "base_model.model.decoder.block.0.layer.0.SelfAttention.q.poly_lora_B", |
| "base_model.model.decoder.block.0.layer.0.SelfAttention.q.poly_router.module_logits", |
| "base_model.model.decoder.block.0.layer.0.SelfAttention.v.poly_lora_A", |
| "base_model.model.decoder.block.0.layer.0.SelfAttention.v.poly_lora_B", |
| "base_model.model.decoder.block.0.layer.0.SelfAttention.v.poly_router.module_logits", |
| "base_model.model.decoder.block.0.layer.1.EncDecAttention.q.poly_lora_A", |
| "base_model.model.decoder.block.0.layer.1.EncDecAttention.q.poly_lora_B", |
| "base_model.model.decoder.block.0.layer.1.EncDecAttention.q.poly_router.module_logits", |
| "base_model.model.decoder.block.0.layer.1.EncDecAttention.v.poly_lora_A", |
| "base_model.model.decoder.block.0.layer.1.EncDecAttention.v.poly_lora_B", |
| "base_model.model.decoder.block.0.layer.1.EncDecAttention.v.poly_router.module_logits", |
| "base_model.model.decoder.block.1.layer.0.SelfAttention.q.poly_lora_A", |
| "base_model.model.decoder.block.1.layer.0.SelfAttention.q.poly_lora_B", |
| "base_model.model.decoder.block.1.layer.0.SelfAttention.q.poly_router.module_logits", |
| "base_model.model.decoder.block.1.layer.0.SelfAttention.v.poly_lora_A", |
| "base_model.model.decoder.block.1.layer.0.SelfAttention.v.poly_lora_B", |
| "base_model.model.decoder.block.1.layer.0.SelfAttention.v.poly_router.module_logits", |
| "base_model.model.decoder.block.1.layer.1.EncDecAttention.q.poly_lora_A", |
| "base_model.model.decoder.block.1.layer.1.EncDecAttention.q.poly_lora_B", |
| "base_model.model.decoder.block.1.layer.1.EncDecAttention.q.poly_router.module_logits", |
| "base_model.model.decoder.block.1.layer.1.EncDecAttention.v.poly_lora_A", |
| "base_model.model.decoder.block.1.layer.1.EncDecAttention.v.poly_lora_B", |
| "base_model.model.decoder.block.1.layer.1.EncDecAttention.v.poly_router.module_logits", |
| "base_model.model.decoder.block.2.layer.0.SelfAttention.q.poly_lora_A", |
| "base_model.model.decoder.block.2.layer.0.SelfAttention.q.poly_lora_B", |
| "base_model.model.decoder.block.2.layer.0.SelfAttention.q.poly_router.module_logits", |
| "base_model.model.decoder.block.2.layer.0.SelfAttention.v.poly_lora_A", |
| "base_model.model.decoder.block.2.layer.0.SelfAttention.v.poly_lora_B", |
| "base_model.model.decoder.block.2.layer.0.SelfAttention.v.poly_router.module_logits", |
| "base_model.model.decoder.block.2.layer.1.EncDecAttention.q.poly_lora_A", |
| "base_model.model.decoder.block.2.layer.1.EncDecAttention.q.poly_lora_B", |
| "base_model.model.decoder.block.2.layer.1.EncDecAttention.q.poly_router.module_logits", |
| "base_model.model.decoder.block.2.layer.1.EncDecAttention.v.poly_lora_A", |
| "base_model.model.decoder.block.2.layer.1.EncDecAttention.v.poly_lora_B", |
| "base_model.model.decoder.block.2.layer.1.EncDecAttention.v.poly_router.module_logits", |
| "base_model.model.decoder.block.3.layer.0.SelfAttention.q.poly_lora_A", |
| "base_model.model.decoder.block.3.layer.0.SelfAttention.q.poly_lora_B", |
| "base_model.model.decoder.block.3.layer.0.SelfAttention.q.poly_router.module_logits", |
| "base_model.model.decoder.block.3.layer.0.SelfAttention.v.poly_lora_A", |
| "base_model.model.decoder.block.3.layer.0.SelfAttention.v.poly_lora_B", |
| "base_model.model.decoder.block.3.layer.0.SelfAttention.v.poly_router.module_logits", |
| "base_model.model.decoder.block.3.layer.1.EncDecAttention.q.poly_lora_A", |
| "base_model.model.decoder.block.3.layer.1.EncDecAttention.q.poly_lora_B", |
| "base_model.model.decoder.block.3.layer.1.EncDecAttention.q.poly_router.module_logits", |
| "base_model.model.decoder.block.3.layer.1.EncDecAttention.v.poly_lora_A", |
| "base_model.model.decoder.block.3.layer.1.EncDecAttention.v.poly_lora_B", |
| "base_model.model.decoder.block.3.layer.1.EncDecAttention.v.poly_router.module_logits", |
| "base_model.model.decoder.block.4.layer.0.SelfAttention.q.poly_lora_A", |
| "base_model.model.decoder.block.4.layer.0.SelfAttention.q.poly_lora_B", |
| "base_model.model.decoder.block.4.layer.0.SelfAttention.q.poly_router.module_logits", |
| "base_model.model.decoder.block.4.layer.0.SelfAttention.v.poly_lora_A", |
| "base_model.model.decoder.block.4.layer.0.SelfAttention.v.poly_lora_B", |
| "base_model.model.decoder.block.4.layer.0.SelfAttention.v.poly_router.module_logits", |
| "base_model.model.decoder.block.4.layer.1.EncDecAttention.q.poly_lora_A", |
| "base_model.model.decoder.block.4.layer.1.EncDecAttention.q.poly_lora_B", |
| "base_model.model.decoder.block.4.layer.1.EncDecAttention.q.poly_router.module_logits", |
| "base_model.model.decoder.block.4.layer.1.EncDecAttention.v.poly_lora_A", |
| "base_model.model.decoder.block.4.layer.1.EncDecAttention.v.poly_lora_B", |
| "base_model.model.decoder.block.4.layer.1.EncDecAttention.v.poly_router.module_logits", |
| "base_model.model.encoder.block.0.layer.0.SelfAttention.q.poly_lora_A", |
| "base_model.model.encoder.block.0.layer.0.SelfAttention.q.poly_lora_B", |
| "base_model.model.encoder.block.0.layer.0.SelfAttention.q.poly_router.module_logits", |
| "base_model.model.encoder.block.0.layer.0.SelfAttention.v.poly_lora_A", |
| "base_model.model.encoder.block.0.layer.0.SelfAttention.v.poly_lora_B", |
| "base_model.model.encoder.block.0.layer.0.SelfAttention.v.poly_router.module_logits", |
| "base_model.model.encoder.block.1.layer.0.SelfAttention.q.poly_lora_A", |
| "base_model.model.encoder.block.1.layer.0.SelfAttention.q.poly_lora_B", |
| "base_model.model.encoder.block.1.layer.0.SelfAttention.q.poly_router.module_logits", |
| "base_model.model.encoder.block.1.layer.0.SelfAttention.v.poly_lora_A", |
| "base_model.model.encoder.block.1.layer.0.SelfAttention.v.poly_lora_B", |
| "base_model.model.encoder.block.1.layer.0.SelfAttention.v.poly_router.module_logits", |
| "base_model.model.encoder.block.2.layer.0.SelfAttention.q.poly_lora_A", |
| "base_model.model.encoder.block.2.layer.0.SelfAttention.q.poly_lora_B", |
| "base_model.model.encoder.block.2.layer.0.SelfAttention.q.poly_router.module_logits", |
| "base_model.model.encoder.block.2.layer.0.SelfAttention.v.poly_lora_A", |
| "base_model.model.encoder.block.2.layer.0.SelfAttention.v.poly_lora_B", |
| "base_model.model.encoder.block.2.layer.0.SelfAttention.v.poly_router.module_logits", |
| "base_model.model.encoder.block.3.layer.0.SelfAttention.q.poly_lora_A", |
| "base_model.model.encoder.block.3.layer.0.SelfAttention.q.poly_lora_B", |
| "base_model.model.encoder.block.3.layer.0.SelfAttention.q.poly_router.module_logits", |
| "base_model.model.encoder.block.3.layer.0.SelfAttention.v.poly_lora_A", |
| "base_model.model.encoder.block.3.layer.0.SelfAttention.v.poly_lora_B", |
| "base_model.model.encoder.block.3.layer.0.SelfAttention.v.poly_router.module_logits", |
| "base_model.model.encoder.block.4.layer.0.SelfAttention.q.poly_lora_A", |
| "base_model.model.encoder.block.4.layer.0.SelfAttention.q.poly_lora_B", |
| "base_model.model.encoder.block.4.layer.0.SelfAttention.q.poly_router.module_logits", |
| "base_model.model.encoder.block.4.layer.0.SelfAttention.v.poly_lora_A", |
| "base_model.model.encoder.block.4.layer.0.SelfAttention.v.poly_lora_B", |
| "base_model.model.encoder.block.4.layer.0.SelfAttention.v.poly_router.module_logits" |
| ], |
| "notes": "" |
| } |