{ "_name_or_path": "./best_model", "architectures": [ "BertForSequenceClassification" ], "attention_probs_dropout_prob": 0.1, "classifier_dropout": null, "gradient_checkpointing": false, "hidden_act": "gelu", "hidden_dropout_prob": 0.1, "hidden_size": 768, "id2label": { "0": "1.OA.A.1", "1": "1.OA.A.2", "2": "1.OA.D.8", "3": "2.MD.B.5", "4": "2.MD.C.8", "5": "2.NBT.B.5", "6": "2.NBT.B.6", "7": "2.NBT.B.7", "8": "2.OA.A.1", "9": "3.MD.D.8", "10": "3.NBT.A.2", "11": "3.OA.A.3", "12": "3.OA.A.4", "13": "3.OA.C.7", "14": "3.OA.D.8", "15": "4.MD.A.2", "16": "4.MD.A.3", "17": "4.NBT.B.4", "18": "4.NBT.B.5", "19": "4.NBT.B.6", "20": "4.NF.A.2", "21": "4.OA.A.3", "22": "4.OA.B.4", "23": "5.NBT.B.5", "24": "5.NBT.B.6", "25": "5.NBT.B.7", "26": "5.NF.A.1", "27": "5.NF.A.2", "28": "5.NF.B.4", "29": "5.OA.A.1", "30": "6.EE.A.1", "31": "6.EE.B.7", "32": "6.NS.B.2", "33": "6.NS.B.3", "34": "7.NS.A.1", "35": "7.NS.A.2", "36": "7.NS.A.3", "37": "8.EE.A.2", "38": "8.EE.C.7", "39": "8.EE.C.8", "40": "K.CC.C.7", "41": "K.NBT.A.1", "42": "K.OA.A.4", "43": "K.OA.A.5" }, "initializer_range": 0.02, "intermediate_size": 3072, "label2id": { "1.OA.A.1": 0, "1.OA.A.2": 1, "1.OA.D.8": 2, "2.MD.B.5": 3, "2.MD.C.8": 4, "2.NBT.B.5": 5, "2.NBT.B.6": 6, "2.NBT.B.7": 7, "2.OA.A.1": 8, "3.MD.D.8": 9, "3.NBT.A.2": 10, "3.OA.A.3": 11, "3.OA.A.4": 12, "3.OA.C.7": 13, "3.OA.D.8": 14, "4.MD.A.2": 15, "4.MD.A.3": 16, "4.NBT.B.4": 17, "4.NBT.B.5": 18, "4.NBT.B.6": 19, "4.NF.A.2": 20, "4.OA.A.3": 21, "4.OA.B.4": 22, "5.NBT.B.5": 23, "5.NBT.B.6": 24, "5.NBT.B.7": 25, "5.NF.A.1": 26, "5.NF.A.2": 27, "5.NF.B.4": 28, "5.OA.A.1": 29, "6.EE.A.1": 30, "6.EE.B.7": 31, "6.NS.B.2": 32, "6.NS.B.3": 33, "7.NS.A.1": 34, "7.NS.A.2": 35, "7.NS.A.3": 36, "8.EE.A.2": 37, "8.EE.C.7": 38, "8.EE.C.8": 39, "K.CC.C.7": 40, "K.NBT.A.1": 41, "K.OA.A.4": 42, "K.OA.A.5": 43 }, "layer_norm_eps": 1e-12, "max_position_embeddings": 512, "model_type": "bert", "num_attention_heads": 12, "num_hidden_layers": 12, "pad_token_id": 0, "position_embedding_type": "absolute", "problem_type": "single_label_classification", "torch_dtype": "float32", "transformers_version": "4.47.1", "type_vocab_size": 2, "use_cache": true, "vocab_size": 30522 }