Buckets:
| attention: | |
| head_dim: 0 | |
| head_size: 0 | |
| num_heads: 0 | |
| attribute_tokens: false | |
| auto_batch_size: false | |
| data: | |
| chunk_length: 0 | |
| completion_column: answer | |
| conversation_column: '' | |
| data_kwargs: '' | |
| dataset: vector_data/role_completions/programmer.json | |
| format_template: '' | |
| prompt_column: question | |
| reward_column: '' | |
| skip_nan_rewards: false | |
| split: train | |
| subset: null | |
| truncation: true | |
| debug: false | |
| distributed: | |
| nnode: 1 | |
| node_rank: null | |
| nproc_per_node: 1 | |
| drop_columns: true | |
| filter_modules: layers.1[7-9].*,layers.2[0-9].*,layers.3[0-1].* | |
| force_math_sdp: false | |
| fsdp: false | |
| include_bias: false | |
| label_smoothing: 0.0 | |
| loss_fn: vector_projection | |
| loss_reduction: sum | |
| max_tokens: null | |
| model: allenai/Olmo-3-7B-Instruct | |
| model_kwargs: '' | |
| modules: [] | |
| optimizer_state_path: '' | |
| overwrite: false | |
| peft_init_kwargs: '' | |
| precision: fp32 | |
| processor_path: '' | |
| profile: false | |
| projection_dim: 16 | |
| projection_type: rademacher | |
| reshape_to_square: false | |
| revision: null | |
| run_path: queries/vector/programmer/query_index | |
| skip_index: false | |
| skip_preconditioners: false | |
| split_attention_modules: [] | |
| stats_sample_size: 10000 | |
| stream_shard_size: 400000 | |
| token_batch_size: 4096 | |
| tokenizer: '' | |
| use_tf32_matmuls: false | |
| vector_layer: 17 | |
| vector_path: pt_vectors/programmer.pt | |
Xet Storage Details
- Size:
- 1.25 kB
- Xet hash:
- 361878d9175c42872713d028a3d4dc62ff29dd657b2b0140d70c51f5f35daf88
·
Xet efficiently stores files, intelligently splitting them into unique chunks and accelerating uploads and downloads. More info.