program(1.0) [buildInfo = dict, tensor>({{"coremlc-component-MIL", "3520.4.1"}, {"coremlc-version", "3520.5.1"}, {"coremltools-component-torch", "2.5.1"}, {"coremltools-source-dialect", "TorchScript"}, {"coremltools-version", "9.0"}})] { func main(tensor attention_mask, tensor input_ids) { tensor var_14 = const()[name = tensor("op_14"), val = tensor(1)]; tensor var_21 = const()[name = tensor("op_21"), val = tensor(-1)]; tensor var_34 = not_equal(x = input_ids, y = var_14)[name = tensor("op_34")]; tensor mask_1_dtype_0 = const()[name = tensor("mask_1_dtype_0"), val = tensor("int32")]; tensor var_36_exclusive_0 = const()[name = tensor("op_36_exclusive_0"), val = tensor(false)]; tensor var_36_reverse_0 = const()[name = tensor("op_36_reverse_0"), val = tensor(false)]; tensor mask_1 = cast(dtype = mask_1_dtype_0, x = var_34)[name = tensor("cast_61")]; tensor var_36 = cumsum(axis = var_14, exclusive = var_36_exclusive_0, reverse = var_36_reverse_0, x = mask_1)[name = tensor("op_36")]; tensor incremental_indices = mul(x = var_36, y = mask_1)[name = tensor("incremental_indices")]; tensor var_42 = const()[name = tensor("op_42"), val = tensor(1)]; tensor position_ids = add(x = incremental_indices, y = var_42)[name = tensor("position_ids")]; tensor buffered_token_type_ids_validate_indices_0 = const()[name = tensor("buffered_token_type_ids_validate_indices_0"), val = tensor(false)]; tensor const_2_to_uint16 = const()[name = tensor("const_2_to_uint16"), val = tensor([[0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0, 0]])]; tensor position_ids_to_uint16_dtype_0 = const()[name = tensor("position_ids_to_uint16_dtype_0"), val = tensor("uint16")]; tensor position_ids_to_uint16 = cast(dtype = position_ids_to_uint16_dtype_0, x = position_ids)[name = tensor("cast_60")]; tensor buffered_token_type_ids_cast_uint16 = gather_along_axis(axis = var_14, indices = position_ids_to_uint16, validate_indices = buffered_token_type_ids_validate_indices_0, x = const_2_to_uint16)[name = tensor("buffered_token_type_ids_cast_uint16")]; tensor inputs_embeds_1_batch_dims_0 = const()[name = tensor("inputs_embeds_1_batch_dims_0"), val = tensor(0)]; tensor inputs_embeds_1_validate_indices_0 = const()[name = tensor("inputs_embeds_1_validate_indices_0"), val = tensor(false)]; tensor text_model_embeddings_word_embeddings_weight_to_fp16 = const()[name = tensor("text_model_embeddings_word_embeddings_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(64)))]; tensor greater_equal_0_y_0 = const()[name = tensor("greater_equal_0_y_0"), val = tensor(0)]; tensor greater_equal_0 = greater_equal(x = input_ids, y = greater_equal_0_y_0)[name = tensor("greater_equal_0")]; tensor slice_by_index_0 = const()[name = tensor("slice_by_index_0"), val = tensor(50265)]; tensor add_0 = add(x = input_ids, y = slice_by_index_0)[name = tensor("add_0")]; tensor select_0 = select(a = input_ids, b = add_0, cond = greater_equal_0)[name = tensor("select_0")]; tensor inputs_embeds_1_cast_fp16_axis_0 = const()[name = tensor("inputs_embeds_1_cast_fp16_axis_0"), val = tensor(0)]; tensor inputs_embeds_1_cast_fp16 = gather(axis = inputs_embeds_1_cast_fp16_axis_0, batch_dims = inputs_embeds_1_batch_dims_0, indices = select_0, validate_indices = inputs_embeds_1_validate_indices_0, x = text_model_embeddings_word_embeddings_weight_to_fp16)[name = tensor("inputs_embeds_1_cast_fp16")]; tensor token_type_embeddings_1_axis_0 = const()[name = tensor("token_type_embeddings_1_axis_0"), val = tensor(0)]; tensor token_type_embeddings_1_batch_dims_0 = const()[name = tensor("token_type_embeddings_1_batch_dims_0"), val = tensor(0)]; tensor token_type_embeddings_1_validate_indices_0 = const()[name = tensor("token_type_embeddings_1_validate_indices_0"), val = tensor(false)]; tensor text_model_embeddings_token_type_embeddings_weight_to_fp16 = const()[name = tensor("text_model_embeddings_token_type_embeddings_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(77207168)))]; tensor token_type_embeddings_1_cast_fp16_cast_uint16 = gather(axis = token_type_embeddings_1_axis_0, batch_dims = token_type_embeddings_1_batch_dims_0, indices = buffered_token_type_ids_cast_uint16, validate_indices = token_type_embeddings_1_validate_indices_0, x = text_model_embeddings_token_type_embeddings_weight_to_fp16)[name = tensor("token_type_embeddings_1_cast_fp16_cast_uint16")]; tensor embeddings_1_cast_fp16 = add(x = inputs_embeds_1_cast_fp16, y = token_type_embeddings_1_cast_fp16_cast_uint16)[name = tensor("embeddings_1_cast_fp16")]; tensor position_embeddings_1_axis_0 = const()[name = tensor("position_embeddings_1_axis_0"), val = tensor(0)]; tensor position_embeddings_1_batch_dims_0 = const()[name = tensor("position_embeddings_1_batch_dims_0"), val = tensor(0)]; tensor position_embeddings_1_validate_indices_0 = const()[name = tensor("position_embeddings_1_validate_indices_0"), val = tensor(false)]; tensor text_model_embeddings_position_embeddings_weight_to_fp16 = const()[name = tensor("text_model_embeddings_position_embeddings_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(77208768)))]; tensor position_embeddings_1_cast_fp16_cast_uint16 = gather(axis = position_embeddings_1_axis_0, batch_dims = position_embeddings_1_batch_dims_0, indices = position_ids_to_uint16, validate_indices = position_embeddings_1_validate_indices_0, x = text_model_embeddings_position_embeddings_weight_to_fp16)[name = tensor("position_embeddings_1_cast_fp16_cast_uint16")]; tensor input_3_cast_fp16 = add(x = embeddings_1_cast_fp16, y = position_embeddings_1_cast_fp16_cast_uint16)[name = tensor("input_3_cast_fp16")]; tensor input_5_axes_0 = const()[name = tensor("input_5_axes_0"), val = tensor([-1])]; tensor text_model_embeddings_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_embeddings_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(77998336)))]; tensor text_model_embeddings_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_embeddings_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(77999936)))]; tensor var_23_to_fp16 = const()[name = tensor("op_23_to_fp16"), val = tensor(0x1p-24)]; tensor input_5_cast_fp16 = layer_norm(axes = input_5_axes_0, beta = text_model_embeddings_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_embeddings_LayerNorm_weight_to_fp16, x = input_3_cast_fp16)[name = tensor("input_5_cast_fp16")]; tensor attention_mask_3_dtype_0 = const()[name = tensor("attention_mask_3_dtype_0"), val = tensor("bool")]; tensor const_13 = const()[name = tensor("const_13"), val = tensor([[[[true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true], [true]]]])]; tensor cast_5_dtype_0 = const()[name = tensor("cast_5_dtype_0"), val = tensor("int8")]; tensor gather_nd_0_batch_dims_0 = const()[name = tensor("gather_nd_0_batch_dims_0"), val = tensor(0)]; tensor gather_nd_0_validate_indices_0 = const()[name = tensor("gather_nd_0_validate_indices_0"), val = tensor(false)]; tensor stack_0_to_uint16 = const()[name = tensor("stack_0_to_uint16"), val = tensor([[[[[0, 0], [0, 1], [0, 2], [0, 3], [0, 4], [0, 5], [0, 6], [0, 7], [0, 8], [0, 9], [0, 10], [0, 11], [0, 12], [0, 13], [0, 14], [0, 15], [0, 16], [0, 17], [0, 18], [0, 19], [0, 20], [0, 21], [0, 22], [0, 23], [0, 24], [0, 25], [0, 26], [0, 27], [0, 28], [0, 29], [0, 30], [0, 31], [0, 32], [0, 33], [0, 34], [0, 35], [0, 36], [0, 37], [0, 38], [0, 39], [0, 40], [0, 41], [0, 42], [0, 43], [0, 44], [0, 45], [0, 46], [0, 47], [0, 48], [0, 49], [0, 50], [0, 51], [0, 52], [0, 53], [0, 54], [0, 55], [0, 56], [0, 57], [0, 58], [0, 59], [0, 60], [0, 61], [0, 62], [0, 63], [0, 64], [0, 65], [0, 66], [0, 67], [0, 68], [0, 69], [0, 70], [0, 71], [0, 72], [0, 73], [0, 74], [0, 75], [0, 76], [0, 77], [0, 78], [0, 79], [0, 80], [0, 81], [0, 82], [0, 83], [0, 84], [0, 85], [0, 86], [0, 87], [0, 88], [0, 89], [0, 90], [0, 91], [0, 92], [0, 93], [0, 94], [0, 95], [0, 96], [0, 97], [0, 98], [0, 99], [0, 100], [0, 101], [0, 102], [0, 103], [0, 104], [0, 105], [0, 106], [0, 107], [0, 108], [0, 109], [0, 110], [0, 111], [0, 112], [0, 113], [0, 114], [0, 115], [0, 116], [0, 117], [0, 118], [0, 119], [0, 120], [0, 121], [0, 122], [0, 123], [0, 124], [0, 125], [0, 126], [0, 127], [0, 128], [0, 129], [0, 130], [0, 131], [0, 132], [0, 133], [0, 134], [0, 135], [0, 136], [0, 137], [0, 138], [0, 139], [0, 140], [0, 141], [0, 142], [0, 143], [0, 144], [0, 145], [0, 146], [0, 147], [0, 148], [0, 149], [0, 150], [0, 151], [0, 152], [0, 153], [0, 154], [0, 155], [0, 156], [0, 157], [0, 158], [0, 159], [0, 160], [0, 161], [0, 162], [0, 163], [0, 164], [0, 165], [0, 166], [0, 167], [0, 168], [0, 169], [0, 170], [0, 171], [0, 172], [0, 173], [0, 174], [0, 175], [0, 176], [0, 177], [0, 178], [0, 179], [0, 180], [0, 181], [0, 182], [0, 183], [0, 184], [0, 185], [0, 186], [0, 187], [0, 188], [0, 189], [0, 190], [0, 191], [0, 192], [0, 193], [0, 194], [0, 195], [0, 196], [0, 197], [0, 198], [0, 199], [0, 200], [0, 201], [0, 202], [0, 203], [0, 204], [0, 205], [0, 206], [0, 207], [0, 208], [0, 209], [0, 210], [0, 211], [0, 212], [0, 213], [0, 214], [0, 215], [0, 216], [0, 217], [0, 218], [0, 219], [0, 220], [0, 221], [0, 222], [0, 223], [0, 224], [0, 225], [0, 226], [0, 227], [0, 228], [0, 229], [0, 230], [0, 231], [0, 232], [0, 233], [0, 234], [0, 235], [0, 236], [0, 237], [0, 238], [0, 239], [0, 240], [0, 241], [0, 242], [0, 243], [0, 244], [0, 245], [0, 246], [0, 247], [0, 248], [0, 249], [0, 250], [0, 251], [0, 252], [0, 253], [0, 254], [0, 255], [0, 256], [0, 257], [0, 258], [0, 259], [0, 260], [0, 261], [0, 262], [0, 263], [0, 264], [0, 265], [0, 266], [0, 267], [0, 268], [0, 269], [0, 270], [0, 271], [0, 272], [0, 273], [0, 274], [0, 275], [0, 276], [0, 277], [0, 278], [0, 279], [0, 280], [0, 281], [0, 282], [0, 283], [0, 284], [0, 285], [0, 286], [0, 287], [0, 288], [0, 289], [0, 290], [0, 291], [0, 292], [0, 293], [0, 294], [0, 295], [0, 296], [0, 297], [0, 298], [0, 299], [0, 300], [0, 301], [0, 302], [0, 303], [0, 304], [0, 305], [0, 306], [0, 307], [0, 308], [0, 309], [0, 310], [0, 311], [0, 312], [0, 313], [0, 314], [0, 315], [0, 316], [0, 317], [0, 318], [0, 319], [0, 320], [0, 321], [0, 322], [0, 323], [0, 324], [0, 325], [0, 326], [0, 327], [0, 328], [0, 329], [0, 330], [0, 331], [0, 332], [0, 333], [0, 334], [0, 335], [0, 336], [0, 337], [0, 338], [0, 339], [0, 340], [0, 341], [0, 342], [0, 343], [0, 344], [0, 345], [0, 346], [0, 347], [0, 348], [0, 349], [0, 350], [0, 351], [0, 352], [0, 353], [0, 354], [0, 355], [0, 356], [0, 357], [0, 358], [0, 359], [0, 360], [0, 361], [0, 362], [0, 363], [0, 364], [0, 365], [0, 366], [0, 367], [0, 368], [0, 369], [0, 370], [0, 371], [0, 372], [0, 373], [0, 374], [0, 375], [0, 376], [0, 377], [0, 378], [0, 379], [0, 380], [0, 381], [0, 382], [0, 383], [0, 384], [0, 385], [0, 386], [0, 387], [0, 388], [0, 389], [0, 390], [0, 391], [0, 392], [0, 393], [0, 394], [0, 395], [0, 396], [0, 397], [0, 398], [0, 399], [0, 400], [0, 401], [0, 402], [0, 403], [0, 404], [0, 405], [0, 406], [0, 407], [0, 408], [0, 409], [0, 410], [0, 411], [0, 412], [0, 413], [0, 414], [0, 415], [0, 416], [0, 417], [0, 418], [0, 419], [0, 420], [0, 421], [0, 422], [0, 423], [0, 424], [0, 425], [0, 426], [0, 427], [0, 428], [0, 429], [0, 430], [0, 431], [0, 432], [0, 433], [0, 434], [0, 435], [0, 436], [0, 437], [0, 438], [0, 439], [0, 440], [0, 441], [0, 442], [0, 443], [0, 444], [0, 445], [0, 446], [0, 447], [0, 448], [0, 449], [0, 450], [0, 451], [0, 452], [0, 453], [0, 454], [0, 455], [0, 456], [0, 457], [0, 458], [0, 459], [0, 460], [0, 461], [0, 462], [0, 463], [0, 464], [0, 465], [0, 466], [0, 467], [0, 468], [0, 469], [0, 470], [0, 471], [0, 472], [0, 473], [0, 474], [0, 475], [0, 476], [0, 477], [0, 478], [0, 479], [0, 480], [0, 481], [0, 482], [0, 483], [0, 484], [0, 485], [0, 486], [0, 487], [0, 488], [0, 489], [0, 490], [0, 491], [0, 492], [0, 493], [0, 494], [0, 495], [0, 496], [0, 497], [0, 498], [0, 499], [0, 500], [0, 501], [0, 502], [0, 503], [0, 504], [0, 505], [0, 506], [0, 507], [0, 508], [0, 509], [0, 510], [0, 511]]]]])]; tensor attention_mask_3 = cast(dtype = attention_mask_3_dtype_0, x = attention_mask)[name = tensor("cast_59")]; tensor cast_5 = cast(dtype = cast_5_dtype_0, x = attention_mask_3)[name = tensor("cast_58")]; tensor gather_nd_0_cast_uint16 = gather_nd(batch_dims = gather_nd_0_batch_dims_0, indices = stack_0_to_uint16, validate_indices = gather_nd_0_validate_indices_0, x = cast_5)[name = tensor("gather_nd_0_cast_uint16")]; tensor var_95_transpose_dtype_0 = const()[name = tensor("op_95_transpose_dtype_0"), val = tensor("bool")]; tensor var_95_transpose = cast(dtype = var_95_transpose_dtype_0, x = gather_nd_0_cast_uint16)[name = tensor("cast_57")]; tensor attention_mask_5 = logical_and(x = const_13, y = var_95_transpose)[name = tensor("attention_mask_5")]; tensor const_14_after_broadcast_to_fp16 = const()[name = tensor("const_14_after_broadcast_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(78001536)))]; tensor var_8_after_broadcast_to_fp16 = const()[name = tensor("op_8_after_broadcast_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(78525888)))]; tensor attention_mask_cast_fp16 = select(a = const_14_after_broadcast_to_fp16, b = var_8_after_broadcast_to_fp16, cond = attention_mask_5)[name = tensor("attention_mask_cast_fp16")]; tensor text_model_encoder_layer_0_attention_self_query_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_0_attention_self_query_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(79050240)))]; tensor text_model_encoder_layer_0_attention_self_query_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_0_attention_self_query_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(80229952)))]; tensor linear_0_cast_fp16 = linear(bias = text_model_encoder_layer_0_attention_self_query_bias_to_fp16, weight = text_model_encoder_layer_0_attention_self_query_weight_to_fp16, x = input_5_cast_fp16)[name = tensor("linear_0_cast_fp16")]; tensor var_140 = const()[name = tensor("op_140"), val = tensor([1, 512, -1, 64])]; tensor var_141_cast_fp16 = reshape(shape = var_140, x = linear_0_cast_fp16)[name = tensor("op_141_cast_fp16")]; tensor text_model_encoder_layer_0_attention_self_key_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_0_attention_self_key_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(80231552)))]; tensor text_model_encoder_layer_0_attention_self_key_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_0_attention_self_key_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(81411264)))]; tensor linear_1_cast_fp16 = linear(bias = text_model_encoder_layer_0_attention_self_key_bias_to_fp16, weight = text_model_encoder_layer_0_attention_self_key_weight_to_fp16, x = input_5_cast_fp16)[name = tensor("linear_1_cast_fp16")]; tensor var_146 = const()[name = tensor("op_146"), val = tensor([1, 512, -1, 64])]; tensor var_147_cast_fp16 = reshape(shape = var_146, x = linear_1_cast_fp16)[name = tensor("op_147_cast_fp16")]; tensor text_model_encoder_layer_0_attention_self_value_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_0_attention_self_value_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(81412864)))]; tensor text_model_encoder_layer_0_attention_self_value_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_0_attention_self_value_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(82592576)))]; tensor linear_2_cast_fp16 = linear(bias = text_model_encoder_layer_0_attention_self_value_bias_to_fp16, weight = text_model_encoder_layer_0_attention_self_value_weight_to_fp16, x = input_5_cast_fp16)[name = tensor("linear_2_cast_fp16")]; tensor var_152 = const()[name = tensor("op_152"), val = tensor([1, 512, -1, 64])]; tensor var_153_cast_fp16 = reshape(shape = var_152, x = linear_2_cast_fp16)[name = tensor("op_153_cast_fp16")]; tensor value_1_perm_0 = const()[name = tensor("value_1_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_156_transpose_x_0 = const()[name = tensor("op_156_transpose_x_0"), val = tensor(false)]; tensor var_156_transpose_y_0 = const()[name = tensor("op_156_transpose_y_0"), val = tensor(false)]; tensor transpose_37_perm_0 = const()[name = tensor("transpose_37_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_38_perm_0 = const()[name = tensor("transpose_38_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_38 = transpose(perm = transpose_38_perm_0, x = var_147_cast_fp16)[name = tensor("transpose_106")]; tensor transpose_37 = transpose(perm = transpose_37_perm_0, x = var_141_cast_fp16)[name = tensor("transpose_107")]; tensor var_156_cast_fp16 = matmul(transpose_x = var_156_transpose_x_0, transpose_y = var_156_transpose_y_0, x = transpose_37, y = transpose_38)[name = tensor("op_156_cast_fp16")]; tensor var_157_to_fp16 = const()[name = tensor("op_157_to_fp16"), val = tensor(0x1p-3)]; tensor attn_weights_1_cast_fp16 = mul(x = var_156_cast_fp16, y = var_157_to_fp16)[name = tensor("attn_weights_1_cast_fp16")]; tensor input_7_cast_fp16 = add(x = attn_weights_1_cast_fp16, y = attention_mask_cast_fp16)[name = tensor("input_7_cast_fp16")]; tensor var_160_cast_fp16 = softmax(axis = var_21, x = input_7_cast_fp16)[name = tensor("op_160_cast_fp16")]; tensor attn_output_1_transpose_x_0 = const()[name = tensor("attn_output_1_transpose_x_0"), val = tensor(false)]; tensor attn_output_1_transpose_y_0 = const()[name = tensor("attn_output_1_transpose_y_0"), val = tensor(false)]; tensor value_1_cast_fp16 = transpose(perm = value_1_perm_0, x = var_153_cast_fp16)[name = tensor("transpose_108")]; tensor attn_output_1_cast_fp16 = matmul(transpose_x = attn_output_1_transpose_x_0, transpose_y = attn_output_1_transpose_y_0, x = var_160_cast_fp16, y = value_1_cast_fp16)[name = tensor("attn_output_1_cast_fp16")]; tensor var_164_perm_0 = const()[name = tensor("op_164_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_166 = const()[name = tensor("op_166"), val = tensor([1, 512, -1])]; tensor var_164_cast_fp16 = transpose(perm = var_164_perm_0, x = attn_output_1_cast_fp16)[name = tensor("transpose_105")]; tensor var_167_cast_fp16 = reshape(shape = var_166, x = var_164_cast_fp16)[name = tensor("op_167_cast_fp16")]; tensor text_model_encoder_layer_0_attention_output_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_0_attention_output_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(82594176)))]; tensor text_model_encoder_layer_0_attention_output_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_0_attention_output_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(83773888)))]; tensor linear_3_cast_fp16 = linear(bias = text_model_encoder_layer_0_attention_output_dense_bias_to_fp16, weight = text_model_encoder_layer_0_attention_output_dense_weight_to_fp16, x = var_167_cast_fp16)[name = tensor("linear_3_cast_fp16")]; tensor input_15_cast_fp16 = add(x = linear_3_cast_fp16, y = input_5_cast_fp16)[name = tensor("input_15_cast_fp16")]; tensor input_17_axes_0 = const()[name = tensor("input_17_axes_0"), val = tensor([-1])]; tensor text_model_encoder_layer_0_attention_output_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_0_attention_output_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(83775488)))]; tensor text_model_encoder_layer_0_attention_output_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_0_attention_output_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(83777088)))]; tensor input_17_cast_fp16 = layer_norm(axes = input_17_axes_0, beta = text_model_encoder_layer_0_attention_output_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_encoder_layer_0_attention_output_LayerNorm_weight_to_fp16, x = input_15_cast_fp16)[name = tensor("input_17_cast_fp16")]; tensor text_model_encoder_layer_0_intermediate_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_0_intermediate_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(83778688)))]; tensor text_model_encoder_layer_0_intermediate_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_0_intermediate_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(88497344)))]; tensor linear_4_cast_fp16 = linear(bias = text_model_encoder_layer_0_intermediate_dense_bias_to_fp16, weight = text_model_encoder_layer_0_intermediate_dense_weight_to_fp16, x = input_17_cast_fp16)[name = tensor("linear_4_cast_fp16")]; tensor input_21_mode_0 = const()[name = tensor("input_21_mode_0"), val = tensor("EXACT")]; tensor input_21_cast_fp16 = gelu(mode = input_21_mode_0, x = linear_4_cast_fp16)[name = tensor("input_21_cast_fp16")]; tensor text_model_encoder_layer_0_output_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_0_output_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(88503552)))]; tensor text_model_encoder_layer_0_output_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_0_output_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(93222208)))]; tensor linear_5_cast_fp16 = linear(bias = text_model_encoder_layer_0_output_dense_bias_to_fp16, weight = text_model_encoder_layer_0_output_dense_weight_to_fp16, x = input_21_cast_fp16)[name = tensor("linear_5_cast_fp16")]; tensor input_25_cast_fp16 = add(x = linear_5_cast_fp16, y = input_17_cast_fp16)[name = tensor("input_25_cast_fp16")]; tensor hidden_states_5_axes_0 = const()[name = tensor("hidden_states_5_axes_0"), val = tensor([-1])]; tensor text_model_encoder_layer_0_output_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_0_output_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(93223808)))]; tensor text_model_encoder_layer_0_output_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_0_output_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(93225408)))]; tensor hidden_states_5_cast_fp16 = layer_norm(axes = hidden_states_5_axes_0, beta = text_model_encoder_layer_0_output_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_encoder_layer_0_output_LayerNorm_weight_to_fp16, x = input_25_cast_fp16)[name = tensor("hidden_states_5_cast_fp16")]; tensor text_model_encoder_layer_1_attention_self_query_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_1_attention_self_query_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(93227008)))]; tensor text_model_encoder_layer_1_attention_self_query_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_1_attention_self_query_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(94406720)))]; tensor linear_6_cast_fp16 = linear(bias = text_model_encoder_layer_1_attention_self_query_bias_to_fp16, weight = text_model_encoder_layer_1_attention_self_query_weight_to_fp16, x = hidden_states_5_cast_fp16)[name = tensor("linear_6_cast_fp16")]; tensor var_209 = const()[name = tensor("op_209"), val = tensor([1, 512, -1, 64])]; tensor var_210_cast_fp16 = reshape(shape = var_209, x = linear_6_cast_fp16)[name = tensor("op_210_cast_fp16")]; tensor text_model_encoder_layer_1_attention_self_key_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_1_attention_self_key_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(94408320)))]; tensor text_model_encoder_layer_1_attention_self_key_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_1_attention_self_key_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(95588032)))]; tensor linear_7_cast_fp16 = linear(bias = text_model_encoder_layer_1_attention_self_key_bias_to_fp16, weight = text_model_encoder_layer_1_attention_self_key_weight_to_fp16, x = hidden_states_5_cast_fp16)[name = tensor("linear_7_cast_fp16")]; tensor var_215 = const()[name = tensor("op_215"), val = tensor([1, 512, -1, 64])]; tensor var_216_cast_fp16 = reshape(shape = var_215, x = linear_7_cast_fp16)[name = tensor("op_216_cast_fp16")]; tensor text_model_encoder_layer_1_attention_self_value_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_1_attention_self_value_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(95589632)))]; tensor text_model_encoder_layer_1_attention_self_value_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_1_attention_self_value_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(96769344)))]; tensor linear_8_cast_fp16 = linear(bias = text_model_encoder_layer_1_attention_self_value_bias_to_fp16, weight = text_model_encoder_layer_1_attention_self_value_weight_to_fp16, x = hidden_states_5_cast_fp16)[name = tensor("linear_8_cast_fp16")]; tensor var_221 = const()[name = tensor("op_221"), val = tensor([1, 512, -1, 64])]; tensor var_222_cast_fp16 = reshape(shape = var_221, x = linear_8_cast_fp16)[name = tensor("op_222_cast_fp16")]; tensor value_3_perm_0 = const()[name = tensor("value_3_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_225_transpose_x_0 = const()[name = tensor("op_225_transpose_x_0"), val = tensor(false)]; tensor var_225_transpose_y_0 = const()[name = tensor("op_225_transpose_y_0"), val = tensor(false)]; tensor transpose_39_perm_0 = const()[name = tensor("transpose_39_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_40_perm_0 = const()[name = tensor("transpose_40_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_40 = transpose(perm = transpose_40_perm_0, x = var_216_cast_fp16)[name = tensor("transpose_102")]; tensor transpose_39 = transpose(perm = transpose_39_perm_0, x = var_210_cast_fp16)[name = tensor("transpose_103")]; tensor var_225_cast_fp16 = matmul(transpose_x = var_225_transpose_x_0, transpose_y = var_225_transpose_y_0, x = transpose_39, y = transpose_40)[name = tensor("op_225_cast_fp16")]; tensor var_226_to_fp16 = const()[name = tensor("op_226_to_fp16"), val = tensor(0x1p-3)]; tensor attn_weights_5_cast_fp16 = mul(x = var_225_cast_fp16, y = var_226_to_fp16)[name = tensor("attn_weights_5_cast_fp16")]; tensor input_27_cast_fp16 = add(x = attn_weights_5_cast_fp16, y = attention_mask_cast_fp16)[name = tensor("input_27_cast_fp16")]; tensor var_229_cast_fp16 = softmax(axis = var_21, x = input_27_cast_fp16)[name = tensor("op_229_cast_fp16")]; tensor attn_output_5_transpose_x_0 = const()[name = tensor("attn_output_5_transpose_x_0"), val = tensor(false)]; tensor attn_output_5_transpose_y_0 = const()[name = tensor("attn_output_5_transpose_y_0"), val = tensor(false)]; tensor value_3_cast_fp16 = transpose(perm = value_3_perm_0, x = var_222_cast_fp16)[name = tensor("transpose_104")]; tensor attn_output_5_cast_fp16 = matmul(transpose_x = attn_output_5_transpose_x_0, transpose_y = attn_output_5_transpose_y_0, x = var_229_cast_fp16, y = value_3_cast_fp16)[name = tensor("attn_output_5_cast_fp16")]; tensor var_233_perm_0 = const()[name = tensor("op_233_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_235 = const()[name = tensor("op_235"), val = tensor([1, 512, -1])]; tensor var_233_cast_fp16 = transpose(perm = var_233_perm_0, x = attn_output_5_cast_fp16)[name = tensor("transpose_101")]; tensor var_236_cast_fp16 = reshape(shape = var_235, x = var_233_cast_fp16)[name = tensor("op_236_cast_fp16")]; tensor text_model_encoder_layer_1_attention_output_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_1_attention_output_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(96770944)))]; tensor text_model_encoder_layer_1_attention_output_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_1_attention_output_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(97950656)))]; tensor linear_9_cast_fp16 = linear(bias = text_model_encoder_layer_1_attention_output_dense_bias_to_fp16, weight = text_model_encoder_layer_1_attention_output_dense_weight_to_fp16, x = var_236_cast_fp16)[name = tensor("linear_9_cast_fp16")]; tensor input_35_cast_fp16 = add(x = linear_9_cast_fp16, y = hidden_states_5_cast_fp16)[name = tensor("input_35_cast_fp16")]; tensor input_37_axes_0 = const()[name = tensor("input_37_axes_0"), val = tensor([-1])]; tensor text_model_encoder_layer_1_attention_output_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_1_attention_output_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(97952256)))]; tensor text_model_encoder_layer_1_attention_output_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_1_attention_output_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(97953856)))]; tensor input_37_cast_fp16 = layer_norm(axes = input_37_axes_0, beta = text_model_encoder_layer_1_attention_output_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_encoder_layer_1_attention_output_LayerNorm_weight_to_fp16, x = input_35_cast_fp16)[name = tensor("input_37_cast_fp16")]; tensor text_model_encoder_layer_1_intermediate_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_1_intermediate_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(97955456)))]; tensor text_model_encoder_layer_1_intermediate_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_1_intermediate_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(102674112)))]; tensor linear_10_cast_fp16 = linear(bias = text_model_encoder_layer_1_intermediate_dense_bias_to_fp16, weight = text_model_encoder_layer_1_intermediate_dense_weight_to_fp16, x = input_37_cast_fp16)[name = tensor("linear_10_cast_fp16")]; tensor input_41_mode_0 = const()[name = tensor("input_41_mode_0"), val = tensor("EXACT")]; tensor input_41_cast_fp16 = gelu(mode = input_41_mode_0, x = linear_10_cast_fp16)[name = tensor("input_41_cast_fp16")]; tensor text_model_encoder_layer_1_output_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_1_output_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(102680320)))]; tensor text_model_encoder_layer_1_output_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_1_output_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(107398976)))]; tensor linear_11_cast_fp16 = linear(bias = text_model_encoder_layer_1_output_dense_bias_to_fp16, weight = text_model_encoder_layer_1_output_dense_weight_to_fp16, x = input_41_cast_fp16)[name = tensor("linear_11_cast_fp16")]; tensor input_45_cast_fp16 = add(x = linear_11_cast_fp16, y = input_37_cast_fp16)[name = tensor("input_45_cast_fp16")]; tensor hidden_states_11_axes_0 = const()[name = tensor("hidden_states_11_axes_0"), val = tensor([-1])]; tensor text_model_encoder_layer_1_output_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_1_output_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(107400576)))]; tensor text_model_encoder_layer_1_output_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_1_output_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(107402176)))]; tensor hidden_states_11_cast_fp16 = layer_norm(axes = hidden_states_11_axes_0, beta = text_model_encoder_layer_1_output_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_encoder_layer_1_output_LayerNorm_weight_to_fp16, x = input_45_cast_fp16)[name = tensor("hidden_states_11_cast_fp16")]; tensor text_model_encoder_layer_2_attention_self_query_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_2_attention_self_query_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(107403776)))]; tensor text_model_encoder_layer_2_attention_self_query_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_2_attention_self_query_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(108583488)))]; tensor linear_12_cast_fp16 = linear(bias = text_model_encoder_layer_2_attention_self_query_bias_to_fp16, weight = text_model_encoder_layer_2_attention_self_query_weight_to_fp16, x = hidden_states_11_cast_fp16)[name = tensor("linear_12_cast_fp16")]; tensor var_278 = const()[name = tensor("op_278"), val = tensor([1, 512, -1, 64])]; tensor var_279_cast_fp16 = reshape(shape = var_278, x = linear_12_cast_fp16)[name = tensor("op_279_cast_fp16")]; tensor text_model_encoder_layer_2_attention_self_key_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_2_attention_self_key_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(108585088)))]; tensor text_model_encoder_layer_2_attention_self_key_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_2_attention_self_key_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(109764800)))]; tensor linear_13_cast_fp16 = linear(bias = text_model_encoder_layer_2_attention_self_key_bias_to_fp16, weight = text_model_encoder_layer_2_attention_self_key_weight_to_fp16, x = hidden_states_11_cast_fp16)[name = tensor("linear_13_cast_fp16")]; tensor var_284 = const()[name = tensor("op_284"), val = tensor([1, 512, -1, 64])]; tensor var_285_cast_fp16 = reshape(shape = var_284, x = linear_13_cast_fp16)[name = tensor("op_285_cast_fp16")]; tensor text_model_encoder_layer_2_attention_self_value_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_2_attention_self_value_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(109766400)))]; tensor text_model_encoder_layer_2_attention_self_value_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_2_attention_self_value_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(110946112)))]; tensor linear_14_cast_fp16 = linear(bias = text_model_encoder_layer_2_attention_self_value_bias_to_fp16, weight = text_model_encoder_layer_2_attention_self_value_weight_to_fp16, x = hidden_states_11_cast_fp16)[name = tensor("linear_14_cast_fp16")]; tensor var_290 = const()[name = tensor("op_290"), val = tensor([1, 512, -1, 64])]; tensor var_291_cast_fp16 = reshape(shape = var_290, x = linear_14_cast_fp16)[name = tensor("op_291_cast_fp16")]; tensor value_5_perm_0 = const()[name = tensor("value_5_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_294_transpose_x_0 = const()[name = tensor("op_294_transpose_x_0"), val = tensor(false)]; tensor var_294_transpose_y_0 = const()[name = tensor("op_294_transpose_y_0"), val = tensor(false)]; tensor transpose_41_perm_0 = const()[name = tensor("transpose_41_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_42_perm_0 = const()[name = tensor("transpose_42_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_42 = transpose(perm = transpose_42_perm_0, x = var_285_cast_fp16)[name = tensor("transpose_98")]; tensor transpose_41 = transpose(perm = transpose_41_perm_0, x = var_279_cast_fp16)[name = tensor("transpose_99")]; tensor var_294_cast_fp16 = matmul(transpose_x = var_294_transpose_x_0, transpose_y = var_294_transpose_y_0, x = transpose_41, y = transpose_42)[name = tensor("op_294_cast_fp16")]; tensor var_295_to_fp16 = const()[name = tensor("op_295_to_fp16"), val = tensor(0x1p-3)]; tensor attn_weights_9_cast_fp16 = mul(x = var_294_cast_fp16, y = var_295_to_fp16)[name = tensor("attn_weights_9_cast_fp16")]; tensor input_47_cast_fp16 = add(x = attn_weights_9_cast_fp16, y = attention_mask_cast_fp16)[name = tensor("input_47_cast_fp16")]; tensor var_298_cast_fp16 = softmax(axis = var_21, x = input_47_cast_fp16)[name = tensor("op_298_cast_fp16")]; tensor attn_output_9_transpose_x_0 = const()[name = tensor("attn_output_9_transpose_x_0"), val = tensor(false)]; tensor attn_output_9_transpose_y_0 = const()[name = tensor("attn_output_9_transpose_y_0"), val = tensor(false)]; tensor value_5_cast_fp16 = transpose(perm = value_5_perm_0, x = var_291_cast_fp16)[name = tensor("transpose_100")]; tensor attn_output_9_cast_fp16 = matmul(transpose_x = attn_output_9_transpose_x_0, transpose_y = attn_output_9_transpose_y_0, x = var_298_cast_fp16, y = value_5_cast_fp16)[name = tensor("attn_output_9_cast_fp16")]; tensor var_302_perm_0 = const()[name = tensor("op_302_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_304 = const()[name = tensor("op_304"), val = tensor([1, 512, -1])]; tensor var_302_cast_fp16 = transpose(perm = var_302_perm_0, x = attn_output_9_cast_fp16)[name = tensor("transpose_97")]; tensor var_305_cast_fp16 = reshape(shape = var_304, x = var_302_cast_fp16)[name = tensor("op_305_cast_fp16")]; tensor text_model_encoder_layer_2_attention_output_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_2_attention_output_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(110947712)))]; tensor text_model_encoder_layer_2_attention_output_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_2_attention_output_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(112127424)))]; tensor linear_15_cast_fp16 = linear(bias = text_model_encoder_layer_2_attention_output_dense_bias_to_fp16, weight = text_model_encoder_layer_2_attention_output_dense_weight_to_fp16, x = var_305_cast_fp16)[name = tensor("linear_15_cast_fp16")]; tensor input_55_cast_fp16 = add(x = linear_15_cast_fp16, y = hidden_states_11_cast_fp16)[name = tensor("input_55_cast_fp16")]; tensor input_57_axes_0 = const()[name = tensor("input_57_axes_0"), val = tensor([-1])]; tensor text_model_encoder_layer_2_attention_output_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_2_attention_output_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(112129024)))]; tensor text_model_encoder_layer_2_attention_output_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_2_attention_output_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(112130624)))]; tensor input_57_cast_fp16 = layer_norm(axes = input_57_axes_0, beta = text_model_encoder_layer_2_attention_output_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_encoder_layer_2_attention_output_LayerNorm_weight_to_fp16, x = input_55_cast_fp16)[name = tensor("input_57_cast_fp16")]; tensor text_model_encoder_layer_2_intermediate_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_2_intermediate_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(112132224)))]; tensor text_model_encoder_layer_2_intermediate_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_2_intermediate_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(116850880)))]; tensor linear_16_cast_fp16 = linear(bias = text_model_encoder_layer_2_intermediate_dense_bias_to_fp16, weight = text_model_encoder_layer_2_intermediate_dense_weight_to_fp16, x = input_57_cast_fp16)[name = tensor("linear_16_cast_fp16")]; tensor input_61_mode_0 = const()[name = tensor("input_61_mode_0"), val = tensor("EXACT")]; tensor input_61_cast_fp16 = gelu(mode = input_61_mode_0, x = linear_16_cast_fp16)[name = tensor("input_61_cast_fp16")]; tensor text_model_encoder_layer_2_output_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_2_output_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(116857088)))]; tensor text_model_encoder_layer_2_output_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_2_output_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(121575744)))]; tensor linear_17_cast_fp16 = linear(bias = text_model_encoder_layer_2_output_dense_bias_to_fp16, weight = text_model_encoder_layer_2_output_dense_weight_to_fp16, x = input_61_cast_fp16)[name = tensor("linear_17_cast_fp16")]; tensor input_65_cast_fp16 = add(x = linear_17_cast_fp16, y = input_57_cast_fp16)[name = tensor("input_65_cast_fp16")]; tensor hidden_states_17_axes_0 = const()[name = tensor("hidden_states_17_axes_0"), val = tensor([-1])]; tensor text_model_encoder_layer_2_output_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_2_output_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(121577344)))]; tensor text_model_encoder_layer_2_output_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_2_output_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(121578944)))]; tensor hidden_states_17_cast_fp16 = layer_norm(axes = hidden_states_17_axes_0, beta = text_model_encoder_layer_2_output_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_encoder_layer_2_output_LayerNorm_weight_to_fp16, x = input_65_cast_fp16)[name = tensor("hidden_states_17_cast_fp16")]; tensor text_model_encoder_layer_3_attention_self_query_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_3_attention_self_query_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(121580544)))]; tensor text_model_encoder_layer_3_attention_self_query_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_3_attention_self_query_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(122760256)))]; tensor linear_18_cast_fp16 = linear(bias = text_model_encoder_layer_3_attention_self_query_bias_to_fp16, weight = text_model_encoder_layer_3_attention_self_query_weight_to_fp16, x = hidden_states_17_cast_fp16)[name = tensor("linear_18_cast_fp16")]; tensor var_347 = const()[name = tensor("op_347"), val = tensor([1, 512, -1, 64])]; tensor var_348_cast_fp16 = reshape(shape = var_347, x = linear_18_cast_fp16)[name = tensor("op_348_cast_fp16")]; tensor text_model_encoder_layer_3_attention_self_key_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_3_attention_self_key_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(122761856)))]; tensor text_model_encoder_layer_3_attention_self_key_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_3_attention_self_key_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(123941568)))]; tensor linear_19_cast_fp16 = linear(bias = text_model_encoder_layer_3_attention_self_key_bias_to_fp16, weight = text_model_encoder_layer_3_attention_self_key_weight_to_fp16, x = hidden_states_17_cast_fp16)[name = tensor("linear_19_cast_fp16")]; tensor var_353 = const()[name = tensor("op_353"), val = tensor([1, 512, -1, 64])]; tensor var_354_cast_fp16 = reshape(shape = var_353, x = linear_19_cast_fp16)[name = tensor("op_354_cast_fp16")]; tensor text_model_encoder_layer_3_attention_self_value_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_3_attention_self_value_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(123943168)))]; tensor text_model_encoder_layer_3_attention_self_value_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_3_attention_self_value_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(125122880)))]; tensor linear_20_cast_fp16 = linear(bias = text_model_encoder_layer_3_attention_self_value_bias_to_fp16, weight = text_model_encoder_layer_3_attention_self_value_weight_to_fp16, x = hidden_states_17_cast_fp16)[name = tensor("linear_20_cast_fp16")]; tensor var_359 = const()[name = tensor("op_359"), val = tensor([1, 512, -1, 64])]; tensor var_360_cast_fp16 = reshape(shape = var_359, x = linear_20_cast_fp16)[name = tensor("op_360_cast_fp16")]; tensor value_7_perm_0 = const()[name = tensor("value_7_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_363_transpose_x_0 = const()[name = tensor("op_363_transpose_x_0"), val = tensor(false)]; tensor var_363_transpose_y_0 = const()[name = tensor("op_363_transpose_y_0"), val = tensor(false)]; tensor transpose_43_perm_0 = const()[name = tensor("transpose_43_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_44_perm_0 = const()[name = tensor("transpose_44_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_44 = transpose(perm = transpose_44_perm_0, x = var_354_cast_fp16)[name = tensor("transpose_94")]; tensor transpose_43 = transpose(perm = transpose_43_perm_0, x = var_348_cast_fp16)[name = tensor("transpose_95")]; tensor var_363_cast_fp16 = matmul(transpose_x = var_363_transpose_x_0, transpose_y = var_363_transpose_y_0, x = transpose_43, y = transpose_44)[name = tensor("op_363_cast_fp16")]; tensor var_364_to_fp16 = const()[name = tensor("op_364_to_fp16"), val = tensor(0x1p-3)]; tensor attn_weights_13_cast_fp16 = mul(x = var_363_cast_fp16, y = var_364_to_fp16)[name = tensor("attn_weights_13_cast_fp16")]; tensor input_67_cast_fp16 = add(x = attn_weights_13_cast_fp16, y = attention_mask_cast_fp16)[name = tensor("input_67_cast_fp16")]; tensor var_367_cast_fp16 = softmax(axis = var_21, x = input_67_cast_fp16)[name = tensor("op_367_cast_fp16")]; tensor attn_output_13_transpose_x_0 = const()[name = tensor("attn_output_13_transpose_x_0"), val = tensor(false)]; tensor attn_output_13_transpose_y_0 = const()[name = tensor("attn_output_13_transpose_y_0"), val = tensor(false)]; tensor value_7_cast_fp16 = transpose(perm = value_7_perm_0, x = var_360_cast_fp16)[name = tensor("transpose_96")]; tensor attn_output_13_cast_fp16 = matmul(transpose_x = attn_output_13_transpose_x_0, transpose_y = attn_output_13_transpose_y_0, x = var_367_cast_fp16, y = value_7_cast_fp16)[name = tensor("attn_output_13_cast_fp16")]; tensor var_371_perm_0 = const()[name = tensor("op_371_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_373 = const()[name = tensor("op_373"), val = tensor([1, 512, -1])]; tensor var_371_cast_fp16 = transpose(perm = var_371_perm_0, x = attn_output_13_cast_fp16)[name = tensor("transpose_93")]; tensor var_374_cast_fp16 = reshape(shape = var_373, x = var_371_cast_fp16)[name = tensor("op_374_cast_fp16")]; tensor text_model_encoder_layer_3_attention_output_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_3_attention_output_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(125124480)))]; tensor text_model_encoder_layer_3_attention_output_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_3_attention_output_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(126304192)))]; tensor linear_21_cast_fp16 = linear(bias = text_model_encoder_layer_3_attention_output_dense_bias_to_fp16, weight = text_model_encoder_layer_3_attention_output_dense_weight_to_fp16, x = var_374_cast_fp16)[name = tensor("linear_21_cast_fp16")]; tensor input_75_cast_fp16 = add(x = linear_21_cast_fp16, y = hidden_states_17_cast_fp16)[name = tensor("input_75_cast_fp16")]; tensor input_77_axes_0 = const()[name = tensor("input_77_axes_0"), val = tensor([-1])]; tensor text_model_encoder_layer_3_attention_output_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_3_attention_output_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(126305792)))]; tensor text_model_encoder_layer_3_attention_output_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_3_attention_output_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(126307392)))]; tensor input_77_cast_fp16 = layer_norm(axes = input_77_axes_0, beta = text_model_encoder_layer_3_attention_output_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_encoder_layer_3_attention_output_LayerNorm_weight_to_fp16, x = input_75_cast_fp16)[name = tensor("input_77_cast_fp16")]; tensor text_model_encoder_layer_3_intermediate_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_3_intermediate_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(126308992)))]; tensor text_model_encoder_layer_3_intermediate_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_3_intermediate_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(131027648)))]; tensor linear_22_cast_fp16 = linear(bias = text_model_encoder_layer_3_intermediate_dense_bias_to_fp16, weight = text_model_encoder_layer_3_intermediate_dense_weight_to_fp16, x = input_77_cast_fp16)[name = tensor("linear_22_cast_fp16")]; tensor input_81_mode_0 = const()[name = tensor("input_81_mode_0"), val = tensor("EXACT")]; tensor input_81_cast_fp16 = gelu(mode = input_81_mode_0, x = linear_22_cast_fp16)[name = tensor("input_81_cast_fp16")]; tensor text_model_encoder_layer_3_output_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_3_output_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(131033856)))]; tensor text_model_encoder_layer_3_output_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_3_output_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(135752512)))]; tensor linear_23_cast_fp16 = linear(bias = text_model_encoder_layer_3_output_dense_bias_to_fp16, weight = text_model_encoder_layer_3_output_dense_weight_to_fp16, x = input_81_cast_fp16)[name = tensor("linear_23_cast_fp16")]; tensor input_85_cast_fp16 = add(x = linear_23_cast_fp16, y = input_77_cast_fp16)[name = tensor("input_85_cast_fp16")]; tensor hidden_states_23_axes_0 = const()[name = tensor("hidden_states_23_axes_0"), val = tensor([-1])]; tensor text_model_encoder_layer_3_output_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_3_output_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(135754112)))]; tensor text_model_encoder_layer_3_output_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_3_output_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(135755712)))]; tensor hidden_states_23_cast_fp16 = layer_norm(axes = hidden_states_23_axes_0, beta = text_model_encoder_layer_3_output_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_encoder_layer_3_output_LayerNorm_weight_to_fp16, x = input_85_cast_fp16)[name = tensor("hidden_states_23_cast_fp16")]; tensor text_model_encoder_layer_4_attention_self_query_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_4_attention_self_query_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(135757312)))]; tensor text_model_encoder_layer_4_attention_self_query_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_4_attention_self_query_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(136937024)))]; tensor linear_24_cast_fp16 = linear(bias = text_model_encoder_layer_4_attention_self_query_bias_to_fp16, weight = text_model_encoder_layer_4_attention_self_query_weight_to_fp16, x = hidden_states_23_cast_fp16)[name = tensor("linear_24_cast_fp16")]; tensor var_416 = const()[name = tensor("op_416"), val = tensor([1, 512, -1, 64])]; tensor var_417_cast_fp16 = reshape(shape = var_416, x = linear_24_cast_fp16)[name = tensor("op_417_cast_fp16")]; tensor text_model_encoder_layer_4_attention_self_key_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_4_attention_self_key_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(136938624)))]; tensor text_model_encoder_layer_4_attention_self_key_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_4_attention_self_key_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(138118336)))]; tensor linear_25_cast_fp16 = linear(bias = text_model_encoder_layer_4_attention_self_key_bias_to_fp16, weight = text_model_encoder_layer_4_attention_self_key_weight_to_fp16, x = hidden_states_23_cast_fp16)[name = tensor("linear_25_cast_fp16")]; tensor var_422 = const()[name = tensor("op_422"), val = tensor([1, 512, -1, 64])]; tensor var_423_cast_fp16 = reshape(shape = var_422, x = linear_25_cast_fp16)[name = tensor("op_423_cast_fp16")]; tensor text_model_encoder_layer_4_attention_self_value_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_4_attention_self_value_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(138119936)))]; tensor text_model_encoder_layer_4_attention_self_value_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_4_attention_self_value_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(139299648)))]; tensor linear_26_cast_fp16 = linear(bias = text_model_encoder_layer_4_attention_self_value_bias_to_fp16, weight = text_model_encoder_layer_4_attention_self_value_weight_to_fp16, x = hidden_states_23_cast_fp16)[name = tensor("linear_26_cast_fp16")]; tensor var_428 = const()[name = tensor("op_428"), val = tensor([1, 512, -1, 64])]; tensor var_429_cast_fp16 = reshape(shape = var_428, x = linear_26_cast_fp16)[name = tensor("op_429_cast_fp16")]; tensor value_9_perm_0 = const()[name = tensor("value_9_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_432_transpose_x_0 = const()[name = tensor("op_432_transpose_x_0"), val = tensor(false)]; tensor var_432_transpose_y_0 = const()[name = tensor("op_432_transpose_y_0"), val = tensor(false)]; tensor transpose_45_perm_0 = const()[name = tensor("transpose_45_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_46_perm_0 = const()[name = tensor("transpose_46_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_46 = transpose(perm = transpose_46_perm_0, x = var_423_cast_fp16)[name = tensor("transpose_90")]; tensor transpose_45 = transpose(perm = transpose_45_perm_0, x = var_417_cast_fp16)[name = tensor("transpose_91")]; tensor var_432_cast_fp16 = matmul(transpose_x = var_432_transpose_x_0, transpose_y = var_432_transpose_y_0, x = transpose_45, y = transpose_46)[name = tensor("op_432_cast_fp16")]; tensor var_433_to_fp16 = const()[name = tensor("op_433_to_fp16"), val = tensor(0x1p-3)]; tensor attn_weights_17_cast_fp16 = mul(x = var_432_cast_fp16, y = var_433_to_fp16)[name = tensor("attn_weights_17_cast_fp16")]; tensor input_87_cast_fp16 = add(x = attn_weights_17_cast_fp16, y = attention_mask_cast_fp16)[name = tensor("input_87_cast_fp16")]; tensor var_436_cast_fp16 = softmax(axis = var_21, x = input_87_cast_fp16)[name = tensor("op_436_cast_fp16")]; tensor attn_output_17_transpose_x_0 = const()[name = tensor("attn_output_17_transpose_x_0"), val = tensor(false)]; tensor attn_output_17_transpose_y_0 = const()[name = tensor("attn_output_17_transpose_y_0"), val = tensor(false)]; tensor value_9_cast_fp16 = transpose(perm = value_9_perm_0, x = var_429_cast_fp16)[name = tensor("transpose_92")]; tensor attn_output_17_cast_fp16 = matmul(transpose_x = attn_output_17_transpose_x_0, transpose_y = attn_output_17_transpose_y_0, x = var_436_cast_fp16, y = value_9_cast_fp16)[name = tensor("attn_output_17_cast_fp16")]; tensor var_440_perm_0 = const()[name = tensor("op_440_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_442 = const()[name = tensor("op_442"), val = tensor([1, 512, -1])]; tensor var_440_cast_fp16 = transpose(perm = var_440_perm_0, x = attn_output_17_cast_fp16)[name = tensor("transpose_89")]; tensor var_443_cast_fp16 = reshape(shape = var_442, x = var_440_cast_fp16)[name = tensor("op_443_cast_fp16")]; tensor text_model_encoder_layer_4_attention_output_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_4_attention_output_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(139301248)))]; tensor text_model_encoder_layer_4_attention_output_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_4_attention_output_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(140480960)))]; tensor linear_27_cast_fp16 = linear(bias = text_model_encoder_layer_4_attention_output_dense_bias_to_fp16, weight = text_model_encoder_layer_4_attention_output_dense_weight_to_fp16, x = var_443_cast_fp16)[name = tensor("linear_27_cast_fp16")]; tensor input_95_cast_fp16 = add(x = linear_27_cast_fp16, y = hidden_states_23_cast_fp16)[name = tensor("input_95_cast_fp16")]; tensor input_97_axes_0 = const()[name = tensor("input_97_axes_0"), val = tensor([-1])]; tensor text_model_encoder_layer_4_attention_output_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_4_attention_output_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(140482560)))]; tensor text_model_encoder_layer_4_attention_output_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_4_attention_output_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(140484160)))]; tensor input_97_cast_fp16 = layer_norm(axes = input_97_axes_0, beta = text_model_encoder_layer_4_attention_output_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_encoder_layer_4_attention_output_LayerNorm_weight_to_fp16, x = input_95_cast_fp16)[name = tensor("input_97_cast_fp16")]; tensor text_model_encoder_layer_4_intermediate_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_4_intermediate_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(140485760)))]; tensor text_model_encoder_layer_4_intermediate_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_4_intermediate_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(145204416)))]; tensor linear_28_cast_fp16 = linear(bias = text_model_encoder_layer_4_intermediate_dense_bias_to_fp16, weight = text_model_encoder_layer_4_intermediate_dense_weight_to_fp16, x = input_97_cast_fp16)[name = tensor("linear_28_cast_fp16")]; tensor input_101_mode_0 = const()[name = tensor("input_101_mode_0"), val = tensor("EXACT")]; tensor input_101_cast_fp16 = gelu(mode = input_101_mode_0, x = linear_28_cast_fp16)[name = tensor("input_101_cast_fp16")]; tensor text_model_encoder_layer_4_output_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_4_output_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(145210624)))]; tensor text_model_encoder_layer_4_output_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_4_output_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(149929280)))]; tensor linear_29_cast_fp16 = linear(bias = text_model_encoder_layer_4_output_dense_bias_to_fp16, weight = text_model_encoder_layer_4_output_dense_weight_to_fp16, x = input_101_cast_fp16)[name = tensor("linear_29_cast_fp16")]; tensor input_105_cast_fp16 = add(x = linear_29_cast_fp16, y = input_97_cast_fp16)[name = tensor("input_105_cast_fp16")]; tensor hidden_states_29_axes_0 = const()[name = tensor("hidden_states_29_axes_0"), val = tensor([-1])]; tensor text_model_encoder_layer_4_output_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_4_output_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(149930880)))]; tensor text_model_encoder_layer_4_output_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_4_output_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(149932480)))]; tensor hidden_states_29_cast_fp16 = layer_norm(axes = hidden_states_29_axes_0, beta = text_model_encoder_layer_4_output_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_encoder_layer_4_output_LayerNorm_weight_to_fp16, x = input_105_cast_fp16)[name = tensor("hidden_states_29_cast_fp16")]; tensor text_model_encoder_layer_5_attention_self_query_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_5_attention_self_query_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(149934080)))]; tensor text_model_encoder_layer_5_attention_self_query_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_5_attention_self_query_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(151113792)))]; tensor linear_30_cast_fp16 = linear(bias = text_model_encoder_layer_5_attention_self_query_bias_to_fp16, weight = text_model_encoder_layer_5_attention_self_query_weight_to_fp16, x = hidden_states_29_cast_fp16)[name = tensor("linear_30_cast_fp16")]; tensor var_485 = const()[name = tensor("op_485"), val = tensor([1, 512, -1, 64])]; tensor var_486_cast_fp16 = reshape(shape = var_485, x = linear_30_cast_fp16)[name = tensor("op_486_cast_fp16")]; tensor text_model_encoder_layer_5_attention_self_key_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_5_attention_self_key_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(151115392)))]; tensor text_model_encoder_layer_5_attention_self_key_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_5_attention_self_key_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(152295104)))]; tensor linear_31_cast_fp16 = linear(bias = text_model_encoder_layer_5_attention_self_key_bias_to_fp16, weight = text_model_encoder_layer_5_attention_self_key_weight_to_fp16, x = hidden_states_29_cast_fp16)[name = tensor("linear_31_cast_fp16")]; tensor var_491 = const()[name = tensor("op_491"), val = tensor([1, 512, -1, 64])]; tensor var_492_cast_fp16 = reshape(shape = var_491, x = linear_31_cast_fp16)[name = tensor("op_492_cast_fp16")]; tensor text_model_encoder_layer_5_attention_self_value_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_5_attention_self_value_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(152296704)))]; tensor text_model_encoder_layer_5_attention_self_value_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_5_attention_self_value_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(153476416)))]; tensor linear_32_cast_fp16 = linear(bias = text_model_encoder_layer_5_attention_self_value_bias_to_fp16, weight = text_model_encoder_layer_5_attention_self_value_weight_to_fp16, x = hidden_states_29_cast_fp16)[name = tensor("linear_32_cast_fp16")]; tensor var_497 = const()[name = tensor("op_497"), val = tensor([1, 512, -1, 64])]; tensor var_498_cast_fp16 = reshape(shape = var_497, x = linear_32_cast_fp16)[name = tensor("op_498_cast_fp16")]; tensor value_11_perm_0 = const()[name = tensor("value_11_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_501_transpose_x_0 = const()[name = tensor("op_501_transpose_x_0"), val = tensor(false)]; tensor var_501_transpose_y_0 = const()[name = tensor("op_501_transpose_y_0"), val = tensor(false)]; tensor transpose_47_perm_0 = const()[name = tensor("transpose_47_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_48_perm_0 = const()[name = tensor("transpose_48_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_48 = transpose(perm = transpose_48_perm_0, x = var_492_cast_fp16)[name = tensor("transpose_86")]; tensor transpose_47 = transpose(perm = transpose_47_perm_0, x = var_486_cast_fp16)[name = tensor("transpose_87")]; tensor var_501_cast_fp16 = matmul(transpose_x = var_501_transpose_x_0, transpose_y = var_501_transpose_y_0, x = transpose_47, y = transpose_48)[name = tensor("op_501_cast_fp16")]; tensor var_502_to_fp16 = const()[name = tensor("op_502_to_fp16"), val = tensor(0x1p-3)]; tensor attn_weights_21_cast_fp16 = mul(x = var_501_cast_fp16, y = var_502_to_fp16)[name = tensor("attn_weights_21_cast_fp16")]; tensor input_107_cast_fp16 = add(x = attn_weights_21_cast_fp16, y = attention_mask_cast_fp16)[name = tensor("input_107_cast_fp16")]; tensor var_505_cast_fp16 = softmax(axis = var_21, x = input_107_cast_fp16)[name = tensor("op_505_cast_fp16")]; tensor attn_output_21_transpose_x_0 = const()[name = tensor("attn_output_21_transpose_x_0"), val = tensor(false)]; tensor attn_output_21_transpose_y_0 = const()[name = tensor("attn_output_21_transpose_y_0"), val = tensor(false)]; tensor value_11_cast_fp16 = transpose(perm = value_11_perm_0, x = var_498_cast_fp16)[name = tensor("transpose_88")]; tensor attn_output_21_cast_fp16 = matmul(transpose_x = attn_output_21_transpose_x_0, transpose_y = attn_output_21_transpose_y_0, x = var_505_cast_fp16, y = value_11_cast_fp16)[name = tensor("attn_output_21_cast_fp16")]; tensor var_509_perm_0 = const()[name = tensor("op_509_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_511 = const()[name = tensor("op_511"), val = tensor([1, 512, -1])]; tensor var_509_cast_fp16 = transpose(perm = var_509_perm_0, x = attn_output_21_cast_fp16)[name = tensor("transpose_85")]; tensor var_512_cast_fp16 = reshape(shape = var_511, x = var_509_cast_fp16)[name = tensor("op_512_cast_fp16")]; tensor text_model_encoder_layer_5_attention_output_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_5_attention_output_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(153478016)))]; tensor text_model_encoder_layer_5_attention_output_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_5_attention_output_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(154657728)))]; tensor linear_33_cast_fp16 = linear(bias = text_model_encoder_layer_5_attention_output_dense_bias_to_fp16, weight = text_model_encoder_layer_5_attention_output_dense_weight_to_fp16, x = var_512_cast_fp16)[name = tensor("linear_33_cast_fp16")]; tensor input_115_cast_fp16 = add(x = linear_33_cast_fp16, y = hidden_states_29_cast_fp16)[name = tensor("input_115_cast_fp16")]; tensor input_117_axes_0 = const()[name = tensor("input_117_axes_0"), val = tensor([-1])]; tensor text_model_encoder_layer_5_attention_output_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_5_attention_output_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(154659328)))]; tensor text_model_encoder_layer_5_attention_output_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_5_attention_output_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(154660928)))]; tensor input_117_cast_fp16 = layer_norm(axes = input_117_axes_0, beta = text_model_encoder_layer_5_attention_output_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_encoder_layer_5_attention_output_LayerNorm_weight_to_fp16, x = input_115_cast_fp16)[name = tensor("input_117_cast_fp16")]; tensor text_model_encoder_layer_5_intermediate_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_5_intermediate_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(154662528)))]; tensor text_model_encoder_layer_5_intermediate_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_5_intermediate_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(159381184)))]; tensor linear_34_cast_fp16 = linear(bias = text_model_encoder_layer_5_intermediate_dense_bias_to_fp16, weight = text_model_encoder_layer_5_intermediate_dense_weight_to_fp16, x = input_117_cast_fp16)[name = tensor("linear_34_cast_fp16")]; tensor input_121_mode_0 = const()[name = tensor("input_121_mode_0"), val = tensor("EXACT")]; tensor input_121_cast_fp16 = gelu(mode = input_121_mode_0, x = linear_34_cast_fp16)[name = tensor("input_121_cast_fp16")]; tensor text_model_encoder_layer_5_output_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_5_output_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(159387392)))]; tensor text_model_encoder_layer_5_output_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_5_output_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(164106048)))]; tensor linear_35_cast_fp16 = linear(bias = text_model_encoder_layer_5_output_dense_bias_to_fp16, weight = text_model_encoder_layer_5_output_dense_weight_to_fp16, x = input_121_cast_fp16)[name = tensor("linear_35_cast_fp16")]; tensor input_125_cast_fp16 = add(x = linear_35_cast_fp16, y = input_117_cast_fp16)[name = tensor("input_125_cast_fp16")]; tensor hidden_states_35_axes_0 = const()[name = tensor("hidden_states_35_axes_0"), val = tensor([-1])]; tensor text_model_encoder_layer_5_output_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_5_output_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(164107648)))]; tensor text_model_encoder_layer_5_output_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_5_output_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(164109248)))]; tensor hidden_states_35_cast_fp16 = layer_norm(axes = hidden_states_35_axes_0, beta = text_model_encoder_layer_5_output_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_encoder_layer_5_output_LayerNorm_weight_to_fp16, x = input_125_cast_fp16)[name = tensor("hidden_states_35_cast_fp16")]; tensor text_model_encoder_layer_6_attention_self_query_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_6_attention_self_query_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(164110848)))]; tensor text_model_encoder_layer_6_attention_self_query_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_6_attention_self_query_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(165290560)))]; tensor linear_36_cast_fp16 = linear(bias = text_model_encoder_layer_6_attention_self_query_bias_to_fp16, weight = text_model_encoder_layer_6_attention_self_query_weight_to_fp16, x = hidden_states_35_cast_fp16)[name = tensor("linear_36_cast_fp16")]; tensor var_554 = const()[name = tensor("op_554"), val = tensor([1, 512, -1, 64])]; tensor var_555_cast_fp16 = reshape(shape = var_554, x = linear_36_cast_fp16)[name = tensor("op_555_cast_fp16")]; tensor text_model_encoder_layer_6_attention_self_key_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_6_attention_self_key_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(165292160)))]; tensor text_model_encoder_layer_6_attention_self_key_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_6_attention_self_key_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(166471872)))]; tensor linear_37_cast_fp16 = linear(bias = text_model_encoder_layer_6_attention_self_key_bias_to_fp16, weight = text_model_encoder_layer_6_attention_self_key_weight_to_fp16, x = hidden_states_35_cast_fp16)[name = tensor("linear_37_cast_fp16")]; tensor var_560 = const()[name = tensor("op_560"), val = tensor([1, 512, -1, 64])]; tensor var_561_cast_fp16 = reshape(shape = var_560, x = linear_37_cast_fp16)[name = tensor("op_561_cast_fp16")]; tensor text_model_encoder_layer_6_attention_self_value_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_6_attention_self_value_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(166473472)))]; tensor text_model_encoder_layer_6_attention_self_value_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_6_attention_self_value_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(167653184)))]; tensor linear_38_cast_fp16 = linear(bias = text_model_encoder_layer_6_attention_self_value_bias_to_fp16, weight = text_model_encoder_layer_6_attention_self_value_weight_to_fp16, x = hidden_states_35_cast_fp16)[name = tensor("linear_38_cast_fp16")]; tensor var_566 = const()[name = tensor("op_566"), val = tensor([1, 512, -1, 64])]; tensor var_567_cast_fp16 = reshape(shape = var_566, x = linear_38_cast_fp16)[name = tensor("op_567_cast_fp16")]; tensor value_13_perm_0 = const()[name = tensor("value_13_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_570_transpose_x_0 = const()[name = tensor("op_570_transpose_x_0"), val = tensor(false)]; tensor var_570_transpose_y_0 = const()[name = tensor("op_570_transpose_y_0"), val = tensor(false)]; tensor transpose_49_perm_0 = const()[name = tensor("transpose_49_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_50_perm_0 = const()[name = tensor("transpose_50_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_50 = transpose(perm = transpose_50_perm_0, x = var_561_cast_fp16)[name = tensor("transpose_82")]; tensor transpose_49 = transpose(perm = transpose_49_perm_0, x = var_555_cast_fp16)[name = tensor("transpose_83")]; tensor var_570_cast_fp16 = matmul(transpose_x = var_570_transpose_x_0, transpose_y = var_570_transpose_y_0, x = transpose_49, y = transpose_50)[name = tensor("op_570_cast_fp16")]; tensor var_571_to_fp16 = const()[name = tensor("op_571_to_fp16"), val = tensor(0x1p-3)]; tensor attn_weights_25_cast_fp16 = mul(x = var_570_cast_fp16, y = var_571_to_fp16)[name = tensor("attn_weights_25_cast_fp16")]; tensor input_127_cast_fp16 = add(x = attn_weights_25_cast_fp16, y = attention_mask_cast_fp16)[name = tensor("input_127_cast_fp16")]; tensor var_574_cast_fp16 = softmax(axis = var_21, x = input_127_cast_fp16)[name = tensor("op_574_cast_fp16")]; tensor attn_output_25_transpose_x_0 = const()[name = tensor("attn_output_25_transpose_x_0"), val = tensor(false)]; tensor attn_output_25_transpose_y_0 = const()[name = tensor("attn_output_25_transpose_y_0"), val = tensor(false)]; tensor value_13_cast_fp16 = transpose(perm = value_13_perm_0, x = var_567_cast_fp16)[name = tensor("transpose_84")]; tensor attn_output_25_cast_fp16 = matmul(transpose_x = attn_output_25_transpose_x_0, transpose_y = attn_output_25_transpose_y_0, x = var_574_cast_fp16, y = value_13_cast_fp16)[name = tensor("attn_output_25_cast_fp16")]; tensor var_578_perm_0 = const()[name = tensor("op_578_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_580 = const()[name = tensor("op_580"), val = tensor([1, 512, -1])]; tensor var_578_cast_fp16 = transpose(perm = var_578_perm_0, x = attn_output_25_cast_fp16)[name = tensor("transpose_81")]; tensor var_581_cast_fp16 = reshape(shape = var_580, x = var_578_cast_fp16)[name = tensor("op_581_cast_fp16")]; tensor text_model_encoder_layer_6_attention_output_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_6_attention_output_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(167654784)))]; tensor text_model_encoder_layer_6_attention_output_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_6_attention_output_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(168834496)))]; tensor linear_39_cast_fp16 = linear(bias = text_model_encoder_layer_6_attention_output_dense_bias_to_fp16, weight = text_model_encoder_layer_6_attention_output_dense_weight_to_fp16, x = var_581_cast_fp16)[name = tensor("linear_39_cast_fp16")]; tensor input_135_cast_fp16 = add(x = linear_39_cast_fp16, y = hidden_states_35_cast_fp16)[name = tensor("input_135_cast_fp16")]; tensor input_137_axes_0 = const()[name = tensor("input_137_axes_0"), val = tensor([-1])]; tensor text_model_encoder_layer_6_attention_output_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_6_attention_output_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(168836096)))]; tensor text_model_encoder_layer_6_attention_output_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_6_attention_output_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(168837696)))]; tensor input_137_cast_fp16 = layer_norm(axes = input_137_axes_0, beta = text_model_encoder_layer_6_attention_output_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_encoder_layer_6_attention_output_LayerNorm_weight_to_fp16, x = input_135_cast_fp16)[name = tensor("input_137_cast_fp16")]; tensor text_model_encoder_layer_6_intermediate_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_6_intermediate_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(168839296)))]; tensor text_model_encoder_layer_6_intermediate_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_6_intermediate_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(173557952)))]; tensor linear_40_cast_fp16 = linear(bias = text_model_encoder_layer_6_intermediate_dense_bias_to_fp16, weight = text_model_encoder_layer_6_intermediate_dense_weight_to_fp16, x = input_137_cast_fp16)[name = tensor("linear_40_cast_fp16")]; tensor input_141_mode_0 = const()[name = tensor("input_141_mode_0"), val = tensor("EXACT")]; tensor input_141_cast_fp16 = gelu(mode = input_141_mode_0, x = linear_40_cast_fp16)[name = tensor("input_141_cast_fp16")]; tensor text_model_encoder_layer_6_output_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_6_output_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(173564160)))]; tensor text_model_encoder_layer_6_output_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_6_output_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(178282816)))]; tensor linear_41_cast_fp16 = linear(bias = text_model_encoder_layer_6_output_dense_bias_to_fp16, weight = text_model_encoder_layer_6_output_dense_weight_to_fp16, x = input_141_cast_fp16)[name = tensor("linear_41_cast_fp16")]; tensor input_145_cast_fp16 = add(x = linear_41_cast_fp16, y = input_137_cast_fp16)[name = tensor("input_145_cast_fp16")]; tensor hidden_states_41_axes_0 = const()[name = tensor("hidden_states_41_axes_0"), val = tensor([-1])]; tensor text_model_encoder_layer_6_output_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_6_output_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(178284416)))]; tensor text_model_encoder_layer_6_output_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_6_output_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(178286016)))]; tensor hidden_states_41_cast_fp16 = layer_norm(axes = hidden_states_41_axes_0, beta = text_model_encoder_layer_6_output_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_encoder_layer_6_output_LayerNorm_weight_to_fp16, x = input_145_cast_fp16)[name = tensor("hidden_states_41_cast_fp16")]; tensor text_model_encoder_layer_7_attention_self_query_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_7_attention_self_query_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(178287616)))]; tensor text_model_encoder_layer_7_attention_self_query_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_7_attention_self_query_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(179467328)))]; tensor linear_42_cast_fp16 = linear(bias = text_model_encoder_layer_7_attention_self_query_bias_to_fp16, weight = text_model_encoder_layer_7_attention_self_query_weight_to_fp16, x = hidden_states_41_cast_fp16)[name = tensor("linear_42_cast_fp16")]; tensor var_623 = const()[name = tensor("op_623"), val = tensor([1, 512, -1, 64])]; tensor var_624_cast_fp16 = reshape(shape = var_623, x = linear_42_cast_fp16)[name = tensor("op_624_cast_fp16")]; tensor text_model_encoder_layer_7_attention_self_key_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_7_attention_self_key_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(179468928)))]; tensor text_model_encoder_layer_7_attention_self_key_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_7_attention_self_key_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(180648640)))]; tensor linear_43_cast_fp16 = linear(bias = text_model_encoder_layer_7_attention_self_key_bias_to_fp16, weight = text_model_encoder_layer_7_attention_self_key_weight_to_fp16, x = hidden_states_41_cast_fp16)[name = tensor("linear_43_cast_fp16")]; tensor var_629 = const()[name = tensor("op_629"), val = tensor([1, 512, -1, 64])]; tensor var_630_cast_fp16 = reshape(shape = var_629, x = linear_43_cast_fp16)[name = tensor("op_630_cast_fp16")]; tensor text_model_encoder_layer_7_attention_self_value_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_7_attention_self_value_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(180650240)))]; tensor text_model_encoder_layer_7_attention_self_value_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_7_attention_self_value_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(181829952)))]; tensor linear_44_cast_fp16 = linear(bias = text_model_encoder_layer_7_attention_self_value_bias_to_fp16, weight = text_model_encoder_layer_7_attention_self_value_weight_to_fp16, x = hidden_states_41_cast_fp16)[name = tensor("linear_44_cast_fp16")]; tensor var_635 = const()[name = tensor("op_635"), val = tensor([1, 512, -1, 64])]; tensor var_636_cast_fp16 = reshape(shape = var_635, x = linear_44_cast_fp16)[name = tensor("op_636_cast_fp16")]; tensor value_15_perm_0 = const()[name = tensor("value_15_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_639_transpose_x_0 = const()[name = tensor("op_639_transpose_x_0"), val = tensor(false)]; tensor var_639_transpose_y_0 = const()[name = tensor("op_639_transpose_y_0"), val = tensor(false)]; tensor transpose_51_perm_0 = const()[name = tensor("transpose_51_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_52_perm_0 = const()[name = tensor("transpose_52_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_52 = transpose(perm = transpose_52_perm_0, x = var_630_cast_fp16)[name = tensor("transpose_78")]; tensor transpose_51 = transpose(perm = transpose_51_perm_0, x = var_624_cast_fp16)[name = tensor("transpose_79")]; tensor var_639_cast_fp16 = matmul(transpose_x = var_639_transpose_x_0, transpose_y = var_639_transpose_y_0, x = transpose_51, y = transpose_52)[name = tensor("op_639_cast_fp16")]; tensor var_640_to_fp16 = const()[name = tensor("op_640_to_fp16"), val = tensor(0x1p-3)]; tensor attn_weights_29_cast_fp16 = mul(x = var_639_cast_fp16, y = var_640_to_fp16)[name = tensor("attn_weights_29_cast_fp16")]; tensor input_147_cast_fp16 = add(x = attn_weights_29_cast_fp16, y = attention_mask_cast_fp16)[name = tensor("input_147_cast_fp16")]; tensor var_643_cast_fp16 = softmax(axis = var_21, x = input_147_cast_fp16)[name = tensor("op_643_cast_fp16")]; tensor attn_output_29_transpose_x_0 = const()[name = tensor("attn_output_29_transpose_x_0"), val = tensor(false)]; tensor attn_output_29_transpose_y_0 = const()[name = tensor("attn_output_29_transpose_y_0"), val = tensor(false)]; tensor value_15_cast_fp16 = transpose(perm = value_15_perm_0, x = var_636_cast_fp16)[name = tensor("transpose_80")]; tensor attn_output_29_cast_fp16 = matmul(transpose_x = attn_output_29_transpose_x_0, transpose_y = attn_output_29_transpose_y_0, x = var_643_cast_fp16, y = value_15_cast_fp16)[name = tensor("attn_output_29_cast_fp16")]; tensor var_647_perm_0 = const()[name = tensor("op_647_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_649 = const()[name = tensor("op_649"), val = tensor([1, 512, -1])]; tensor var_647_cast_fp16 = transpose(perm = var_647_perm_0, x = attn_output_29_cast_fp16)[name = tensor("transpose_77")]; tensor var_650_cast_fp16 = reshape(shape = var_649, x = var_647_cast_fp16)[name = tensor("op_650_cast_fp16")]; tensor text_model_encoder_layer_7_attention_output_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_7_attention_output_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(181831552)))]; tensor text_model_encoder_layer_7_attention_output_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_7_attention_output_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(183011264)))]; tensor linear_45_cast_fp16 = linear(bias = text_model_encoder_layer_7_attention_output_dense_bias_to_fp16, weight = text_model_encoder_layer_7_attention_output_dense_weight_to_fp16, x = var_650_cast_fp16)[name = tensor("linear_45_cast_fp16")]; tensor input_155_cast_fp16 = add(x = linear_45_cast_fp16, y = hidden_states_41_cast_fp16)[name = tensor("input_155_cast_fp16")]; tensor input_157_axes_0 = const()[name = tensor("input_157_axes_0"), val = tensor([-1])]; tensor text_model_encoder_layer_7_attention_output_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_7_attention_output_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(183012864)))]; tensor text_model_encoder_layer_7_attention_output_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_7_attention_output_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(183014464)))]; tensor input_157_cast_fp16 = layer_norm(axes = input_157_axes_0, beta = text_model_encoder_layer_7_attention_output_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_encoder_layer_7_attention_output_LayerNorm_weight_to_fp16, x = input_155_cast_fp16)[name = tensor("input_157_cast_fp16")]; tensor text_model_encoder_layer_7_intermediate_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_7_intermediate_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(183016064)))]; tensor text_model_encoder_layer_7_intermediate_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_7_intermediate_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(187734720)))]; tensor linear_46_cast_fp16 = linear(bias = text_model_encoder_layer_7_intermediate_dense_bias_to_fp16, weight = text_model_encoder_layer_7_intermediate_dense_weight_to_fp16, x = input_157_cast_fp16)[name = tensor("linear_46_cast_fp16")]; tensor input_161_mode_0 = const()[name = tensor("input_161_mode_0"), val = tensor("EXACT")]; tensor input_161_cast_fp16 = gelu(mode = input_161_mode_0, x = linear_46_cast_fp16)[name = tensor("input_161_cast_fp16")]; tensor text_model_encoder_layer_7_output_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_7_output_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(187740928)))]; tensor text_model_encoder_layer_7_output_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_7_output_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(192459584)))]; tensor linear_47_cast_fp16 = linear(bias = text_model_encoder_layer_7_output_dense_bias_to_fp16, weight = text_model_encoder_layer_7_output_dense_weight_to_fp16, x = input_161_cast_fp16)[name = tensor("linear_47_cast_fp16")]; tensor input_165_cast_fp16 = add(x = linear_47_cast_fp16, y = input_157_cast_fp16)[name = tensor("input_165_cast_fp16")]; tensor hidden_states_47_axes_0 = const()[name = tensor("hidden_states_47_axes_0"), val = tensor([-1])]; tensor text_model_encoder_layer_7_output_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_7_output_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(192461184)))]; tensor text_model_encoder_layer_7_output_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_7_output_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(192462784)))]; tensor hidden_states_47_cast_fp16 = layer_norm(axes = hidden_states_47_axes_0, beta = text_model_encoder_layer_7_output_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_encoder_layer_7_output_LayerNorm_weight_to_fp16, x = input_165_cast_fp16)[name = tensor("hidden_states_47_cast_fp16")]; tensor text_model_encoder_layer_8_attention_self_query_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_8_attention_self_query_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(192464384)))]; tensor text_model_encoder_layer_8_attention_self_query_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_8_attention_self_query_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(193644096)))]; tensor linear_48_cast_fp16 = linear(bias = text_model_encoder_layer_8_attention_self_query_bias_to_fp16, weight = text_model_encoder_layer_8_attention_self_query_weight_to_fp16, x = hidden_states_47_cast_fp16)[name = tensor("linear_48_cast_fp16")]; tensor var_692 = const()[name = tensor("op_692"), val = tensor([1, 512, -1, 64])]; tensor var_693_cast_fp16 = reshape(shape = var_692, x = linear_48_cast_fp16)[name = tensor("op_693_cast_fp16")]; tensor text_model_encoder_layer_8_attention_self_key_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_8_attention_self_key_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(193645696)))]; tensor text_model_encoder_layer_8_attention_self_key_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_8_attention_self_key_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(194825408)))]; tensor linear_49_cast_fp16 = linear(bias = text_model_encoder_layer_8_attention_self_key_bias_to_fp16, weight = text_model_encoder_layer_8_attention_self_key_weight_to_fp16, x = hidden_states_47_cast_fp16)[name = tensor("linear_49_cast_fp16")]; tensor var_698 = const()[name = tensor("op_698"), val = tensor([1, 512, -1, 64])]; tensor var_699_cast_fp16 = reshape(shape = var_698, x = linear_49_cast_fp16)[name = tensor("op_699_cast_fp16")]; tensor text_model_encoder_layer_8_attention_self_value_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_8_attention_self_value_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(194827008)))]; tensor text_model_encoder_layer_8_attention_self_value_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_8_attention_self_value_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(196006720)))]; tensor linear_50_cast_fp16 = linear(bias = text_model_encoder_layer_8_attention_self_value_bias_to_fp16, weight = text_model_encoder_layer_8_attention_self_value_weight_to_fp16, x = hidden_states_47_cast_fp16)[name = tensor("linear_50_cast_fp16")]; tensor var_704 = const()[name = tensor("op_704"), val = tensor([1, 512, -1, 64])]; tensor var_705_cast_fp16 = reshape(shape = var_704, x = linear_50_cast_fp16)[name = tensor("op_705_cast_fp16")]; tensor value_17_perm_0 = const()[name = tensor("value_17_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_708_transpose_x_0 = const()[name = tensor("op_708_transpose_x_0"), val = tensor(false)]; tensor var_708_transpose_y_0 = const()[name = tensor("op_708_transpose_y_0"), val = tensor(false)]; tensor transpose_53_perm_0 = const()[name = tensor("transpose_53_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_54_perm_0 = const()[name = tensor("transpose_54_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_54 = transpose(perm = transpose_54_perm_0, x = var_699_cast_fp16)[name = tensor("transpose_74")]; tensor transpose_53 = transpose(perm = transpose_53_perm_0, x = var_693_cast_fp16)[name = tensor("transpose_75")]; tensor var_708_cast_fp16 = matmul(transpose_x = var_708_transpose_x_0, transpose_y = var_708_transpose_y_0, x = transpose_53, y = transpose_54)[name = tensor("op_708_cast_fp16")]; tensor var_709_to_fp16 = const()[name = tensor("op_709_to_fp16"), val = tensor(0x1p-3)]; tensor attn_weights_33_cast_fp16 = mul(x = var_708_cast_fp16, y = var_709_to_fp16)[name = tensor("attn_weights_33_cast_fp16")]; tensor input_167_cast_fp16 = add(x = attn_weights_33_cast_fp16, y = attention_mask_cast_fp16)[name = tensor("input_167_cast_fp16")]; tensor var_712_cast_fp16 = softmax(axis = var_21, x = input_167_cast_fp16)[name = tensor("op_712_cast_fp16")]; tensor attn_output_33_transpose_x_0 = const()[name = tensor("attn_output_33_transpose_x_0"), val = tensor(false)]; tensor attn_output_33_transpose_y_0 = const()[name = tensor("attn_output_33_transpose_y_0"), val = tensor(false)]; tensor value_17_cast_fp16 = transpose(perm = value_17_perm_0, x = var_705_cast_fp16)[name = tensor("transpose_76")]; tensor attn_output_33_cast_fp16 = matmul(transpose_x = attn_output_33_transpose_x_0, transpose_y = attn_output_33_transpose_y_0, x = var_712_cast_fp16, y = value_17_cast_fp16)[name = tensor("attn_output_33_cast_fp16")]; tensor var_716_perm_0 = const()[name = tensor("op_716_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_718 = const()[name = tensor("op_718"), val = tensor([1, 512, -1])]; tensor var_716_cast_fp16 = transpose(perm = var_716_perm_0, x = attn_output_33_cast_fp16)[name = tensor("transpose_73")]; tensor var_719_cast_fp16 = reshape(shape = var_718, x = var_716_cast_fp16)[name = tensor("op_719_cast_fp16")]; tensor text_model_encoder_layer_8_attention_output_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_8_attention_output_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(196008320)))]; tensor text_model_encoder_layer_8_attention_output_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_8_attention_output_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(197188032)))]; tensor linear_51_cast_fp16 = linear(bias = text_model_encoder_layer_8_attention_output_dense_bias_to_fp16, weight = text_model_encoder_layer_8_attention_output_dense_weight_to_fp16, x = var_719_cast_fp16)[name = tensor("linear_51_cast_fp16")]; tensor input_175_cast_fp16 = add(x = linear_51_cast_fp16, y = hidden_states_47_cast_fp16)[name = tensor("input_175_cast_fp16")]; tensor input_177_axes_0 = const()[name = tensor("input_177_axes_0"), val = tensor([-1])]; tensor text_model_encoder_layer_8_attention_output_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_8_attention_output_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(197189632)))]; tensor text_model_encoder_layer_8_attention_output_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_8_attention_output_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(197191232)))]; tensor input_177_cast_fp16 = layer_norm(axes = input_177_axes_0, beta = text_model_encoder_layer_8_attention_output_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_encoder_layer_8_attention_output_LayerNorm_weight_to_fp16, x = input_175_cast_fp16)[name = tensor("input_177_cast_fp16")]; tensor text_model_encoder_layer_8_intermediate_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_8_intermediate_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(197192832)))]; tensor text_model_encoder_layer_8_intermediate_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_8_intermediate_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(201911488)))]; tensor linear_52_cast_fp16 = linear(bias = text_model_encoder_layer_8_intermediate_dense_bias_to_fp16, weight = text_model_encoder_layer_8_intermediate_dense_weight_to_fp16, x = input_177_cast_fp16)[name = tensor("linear_52_cast_fp16")]; tensor input_181_mode_0 = const()[name = tensor("input_181_mode_0"), val = tensor("EXACT")]; tensor input_181_cast_fp16 = gelu(mode = input_181_mode_0, x = linear_52_cast_fp16)[name = tensor("input_181_cast_fp16")]; tensor text_model_encoder_layer_8_output_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_8_output_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(201917696)))]; tensor text_model_encoder_layer_8_output_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_8_output_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(206636352)))]; tensor linear_53_cast_fp16 = linear(bias = text_model_encoder_layer_8_output_dense_bias_to_fp16, weight = text_model_encoder_layer_8_output_dense_weight_to_fp16, x = input_181_cast_fp16)[name = tensor("linear_53_cast_fp16")]; tensor input_185_cast_fp16 = add(x = linear_53_cast_fp16, y = input_177_cast_fp16)[name = tensor("input_185_cast_fp16")]; tensor hidden_states_53_axes_0 = const()[name = tensor("hidden_states_53_axes_0"), val = tensor([-1])]; tensor text_model_encoder_layer_8_output_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_8_output_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(206637952)))]; tensor text_model_encoder_layer_8_output_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_8_output_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(206639552)))]; tensor hidden_states_53_cast_fp16 = layer_norm(axes = hidden_states_53_axes_0, beta = text_model_encoder_layer_8_output_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_encoder_layer_8_output_LayerNorm_weight_to_fp16, x = input_185_cast_fp16)[name = tensor("hidden_states_53_cast_fp16")]; tensor text_model_encoder_layer_9_attention_self_query_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_9_attention_self_query_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(206641152)))]; tensor text_model_encoder_layer_9_attention_self_query_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_9_attention_self_query_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(207820864)))]; tensor linear_54_cast_fp16 = linear(bias = text_model_encoder_layer_9_attention_self_query_bias_to_fp16, weight = text_model_encoder_layer_9_attention_self_query_weight_to_fp16, x = hidden_states_53_cast_fp16)[name = tensor("linear_54_cast_fp16")]; tensor var_761 = const()[name = tensor("op_761"), val = tensor([1, 512, -1, 64])]; tensor var_762_cast_fp16 = reshape(shape = var_761, x = linear_54_cast_fp16)[name = tensor("op_762_cast_fp16")]; tensor text_model_encoder_layer_9_attention_self_key_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_9_attention_self_key_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(207822464)))]; tensor text_model_encoder_layer_9_attention_self_key_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_9_attention_self_key_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(209002176)))]; tensor linear_55_cast_fp16 = linear(bias = text_model_encoder_layer_9_attention_self_key_bias_to_fp16, weight = text_model_encoder_layer_9_attention_self_key_weight_to_fp16, x = hidden_states_53_cast_fp16)[name = tensor("linear_55_cast_fp16")]; tensor var_767 = const()[name = tensor("op_767"), val = tensor([1, 512, -1, 64])]; tensor var_768_cast_fp16 = reshape(shape = var_767, x = linear_55_cast_fp16)[name = tensor("op_768_cast_fp16")]; tensor text_model_encoder_layer_9_attention_self_value_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_9_attention_self_value_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(209003776)))]; tensor text_model_encoder_layer_9_attention_self_value_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_9_attention_self_value_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(210183488)))]; tensor linear_56_cast_fp16 = linear(bias = text_model_encoder_layer_9_attention_self_value_bias_to_fp16, weight = text_model_encoder_layer_9_attention_self_value_weight_to_fp16, x = hidden_states_53_cast_fp16)[name = tensor("linear_56_cast_fp16")]; tensor var_773 = const()[name = tensor("op_773"), val = tensor([1, 512, -1, 64])]; tensor var_774_cast_fp16 = reshape(shape = var_773, x = linear_56_cast_fp16)[name = tensor("op_774_cast_fp16")]; tensor value_19_perm_0 = const()[name = tensor("value_19_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_777_transpose_x_0 = const()[name = tensor("op_777_transpose_x_0"), val = tensor(false)]; tensor var_777_transpose_y_0 = const()[name = tensor("op_777_transpose_y_0"), val = tensor(false)]; tensor transpose_55_perm_0 = const()[name = tensor("transpose_55_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_56_perm_0 = const()[name = tensor("transpose_56_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_56 = transpose(perm = transpose_56_perm_0, x = var_768_cast_fp16)[name = tensor("transpose_70")]; tensor transpose_55 = transpose(perm = transpose_55_perm_0, x = var_762_cast_fp16)[name = tensor("transpose_71")]; tensor var_777_cast_fp16 = matmul(transpose_x = var_777_transpose_x_0, transpose_y = var_777_transpose_y_0, x = transpose_55, y = transpose_56)[name = tensor("op_777_cast_fp16")]; tensor var_778_to_fp16 = const()[name = tensor("op_778_to_fp16"), val = tensor(0x1p-3)]; tensor attn_weights_37_cast_fp16 = mul(x = var_777_cast_fp16, y = var_778_to_fp16)[name = tensor("attn_weights_37_cast_fp16")]; tensor input_187_cast_fp16 = add(x = attn_weights_37_cast_fp16, y = attention_mask_cast_fp16)[name = tensor("input_187_cast_fp16")]; tensor var_781_cast_fp16 = softmax(axis = var_21, x = input_187_cast_fp16)[name = tensor("op_781_cast_fp16")]; tensor attn_output_37_transpose_x_0 = const()[name = tensor("attn_output_37_transpose_x_0"), val = tensor(false)]; tensor attn_output_37_transpose_y_0 = const()[name = tensor("attn_output_37_transpose_y_0"), val = tensor(false)]; tensor value_19_cast_fp16 = transpose(perm = value_19_perm_0, x = var_774_cast_fp16)[name = tensor("transpose_72")]; tensor attn_output_37_cast_fp16 = matmul(transpose_x = attn_output_37_transpose_x_0, transpose_y = attn_output_37_transpose_y_0, x = var_781_cast_fp16, y = value_19_cast_fp16)[name = tensor("attn_output_37_cast_fp16")]; tensor var_785_perm_0 = const()[name = tensor("op_785_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_787 = const()[name = tensor("op_787"), val = tensor([1, 512, -1])]; tensor var_785_cast_fp16 = transpose(perm = var_785_perm_0, x = attn_output_37_cast_fp16)[name = tensor("transpose_69")]; tensor var_788_cast_fp16 = reshape(shape = var_787, x = var_785_cast_fp16)[name = tensor("op_788_cast_fp16")]; tensor text_model_encoder_layer_9_attention_output_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_9_attention_output_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(210185088)))]; tensor text_model_encoder_layer_9_attention_output_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_9_attention_output_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(211364800)))]; tensor linear_57_cast_fp16 = linear(bias = text_model_encoder_layer_9_attention_output_dense_bias_to_fp16, weight = text_model_encoder_layer_9_attention_output_dense_weight_to_fp16, x = var_788_cast_fp16)[name = tensor("linear_57_cast_fp16")]; tensor input_195_cast_fp16 = add(x = linear_57_cast_fp16, y = hidden_states_53_cast_fp16)[name = tensor("input_195_cast_fp16")]; tensor input_197_axes_0 = const()[name = tensor("input_197_axes_0"), val = tensor([-1])]; tensor text_model_encoder_layer_9_attention_output_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_9_attention_output_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(211366400)))]; tensor text_model_encoder_layer_9_attention_output_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_9_attention_output_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(211368000)))]; tensor input_197_cast_fp16 = layer_norm(axes = input_197_axes_0, beta = text_model_encoder_layer_9_attention_output_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_encoder_layer_9_attention_output_LayerNorm_weight_to_fp16, x = input_195_cast_fp16)[name = tensor("input_197_cast_fp16")]; tensor text_model_encoder_layer_9_intermediate_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_9_intermediate_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(211369600)))]; tensor text_model_encoder_layer_9_intermediate_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_9_intermediate_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(216088256)))]; tensor linear_58_cast_fp16 = linear(bias = text_model_encoder_layer_9_intermediate_dense_bias_to_fp16, weight = text_model_encoder_layer_9_intermediate_dense_weight_to_fp16, x = input_197_cast_fp16)[name = tensor("linear_58_cast_fp16")]; tensor input_201_mode_0 = const()[name = tensor("input_201_mode_0"), val = tensor("EXACT")]; tensor input_201_cast_fp16 = gelu(mode = input_201_mode_0, x = linear_58_cast_fp16)[name = tensor("input_201_cast_fp16")]; tensor text_model_encoder_layer_9_output_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_9_output_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(216094464)))]; tensor text_model_encoder_layer_9_output_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_9_output_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(220813120)))]; tensor linear_59_cast_fp16 = linear(bias = text_model_encoder_layer_9_output_dense_bias_to_fp16, weight = text_model_encoder_layer_9_output_dense_weight_to_fp16, x = input_201_cast_fp16)[name = tensor("linear_59_cast_fp16")]; tensor input_205_cast_fp16 = add(x = linear_59_cast_fp16, y = input_197_cast_fp16)[name = tensor("input_205_cast_fp16")]; tensor hidden_states_59_axes_0 = const()[name = tensor("hidden_states_59_axes_0"), val = tensor([-1])]; tensor text_model_encoder_layer_9_output_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_9_output_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(220814720)))]; tensor text_model_encoder_layer_9_output_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_9_output_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(220816320)))]; tensor hidden_states_59_cast_fp16 = layer_norm(axes = hidden_states_59_axes_0, beta = text_model_encoder_layer_9_output_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_encoder_layer_9_output_LayerNorm_weight_to_fp16, x = input_205_cast_fp16)[name = tensor("hidden_states_59_cast_fp16")]; tensor text_model_encoder_layer_10_attention_self_query_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_10_attention_self_query_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(220817920)))]; tensor text_model_encoder_layer_10_attention_self_query_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_10_attention_self_query_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(221997632)))]; tensor linear_60_cast_fp16 = linear(bias = text_model_encoder_layer_10_attention_self_query_bias_to_fp16, weight = text_model_encoder_layer_10_attention_self_query_weight_to_fp16, x = hidden_states_59_cast_fp16)[name = tensor("linear_60_cast_fp16")]; tensor var_830 = const()[name = tensor("op_830"), val = tensor([1, 512, -1, 64])]; tensor var_831_cast_fp16 = reshape(shape = var_830, x = linear_60_cast_fp16)[name = tensor("op_831_cast_fp16")]; tensor text_model_encoder_layer_10_attention_self_key_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_10_attention_self_key_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(221999232)))]; tensor text_model_encoder_layer_10_attention_self_key_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_10_attention_self_key_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(223178944)))]; tensor linear_61_cast_fp16 = linear(bias = text_model_encoder_layer_10_attention_self_key_bias_to_fp16, weight = text_model_encoder_layer_10_attention_self_key_weight_to_fp16, x = hidden_states_59_cast_fp16)[name = tensor("linear_61_cast_fp16")]; tensor var_836 = const()[name = tensor("op_836"), val = tensor([1, 512, -1, 64])]; tensor var_837_cast_fp16 = reshape(shape = var_836, x = linear_61_cast_fp16)[name = tensor("op_837_cast_fp16")]; tensor text_model_encoder_layer_10_attention_self_value_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_10_attention_self_value_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(223180544)))]; tensor text_model_encoder_layer_10_attention_self_value_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_10_attention_self_value_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(224360256)))]; tensor linear_62_cast_fp16 = linear(bias = text_model_encoder_layer_10_attention_self_value_bias_to_fp16, weight = text_model_encoder_layer_10_attention_self_value_weight_to_fp16, x = hidden_states_59_cast_fp16)[name = tensor("linear_62_cast_fp16")]; tensor var_842 = const()[name = tensor("op_842"), val = tensor([1, 512, -1, 64])]; tensor var_843_cast_fp16 = reshape(shape = var_842, x = linear_62_cast_fp16)[name = tensor("op_843_cast_fp16")]; tensor value_21_perm_0 = const()[name = tensor("value_21_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_846_transpose_x_0 = const()[name = tensor("op_846_transpose_x_0"), val = tensor(false)]; tensor var_846_transpose_y_0 = const()[name = tensor("op_846_transpose_y_0"), val = tensor(false)]; tensor transpose_57_perm_0 = const()[name = tensor("transpose_57_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_58_perm_0 = const()[name = tensor("transpose_58_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_58 = transpose(perm = transpose_58_perm_0, x = var_837_cast_fp16)[name = tensor("transpose_66")]; tensor transpose_57 = transpose(perm = transpose_57_perm_0, x = var_831_cast_fp16)[name = tensor("transpose_67")]; tensor var_846_cast_fp16 = matmul(transpose_x = var_846_transpose_x_0, transpose_y = var_846_transpose_y_0, x = transpose_57, y = transpose_58)[name = tensor("op_846_cast_fp16")]; tensor var_847_to_fp16 = const()[name = tensor("op_847_to_fp16"), val = tensor(0x1p-3)]; tensor attn_weights_41_cast_fp16 = mul(x = var_846_cast_fp16, y = var_847_to_fp16)[name = tensor("attn_weights_41_cast_fp16")]; tensor input_207_cast_fp16 = add(x = attn_weights_41_cast_fp16, y = attention_mask_cast_fp16)[name = tensor("input_207_cast_fp16")]; tensor var_850_cast_fp16 = softmax(axis = var_21, x = input_207_cast_fp16)[name = tensor("op_850_cast_fp16")]; tensor attn_output_41_transpose_x_0 = const()[name = tensor("attn_output_41_transpose_x_0"), val = tensor(false)]; tensor attn_output_41_transpose_y_0 = const()[name = tensor("attn_output_41_transpose_y_0"), val = tensor(false)]; tensor value_21_cast_fp16 = transpose(perm = value_21_perm_0, x = var_843_cast_fp16)[name = tensor("transpose_68")]; tensor attn_output_41_cast_fp16 = matmul(transpose_x = attn_output_41_transpose_x_0, transpose_y = attn_output_41_transpose_y_0, x = var_850_cast_fp16, y = value_21_cast_fp16)[name = tensor("attn_output_41_cast_fp16")]; tensor var_854_perm_0 = const()[name = tensor("op_854_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_856 = const()[name = tensor("op_856"), val = tensor([1, 512, -1])]; tensor var_854_cast_fp16 = transpose(perm = var_854_perm_0, x = attn_output_41_cast_fp16)[name = tensor("transpose_65")]; tensor var_857_cast_fp16 = reshape(shape = var_856, x = var_854_cast_fp16)[name = tensor("op_857_cast_fp16")]; tensor text_model_encoder_layer_10_attention_output_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_10_attention_output_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(224361856)))]; tensor text_model_encoder_layer_10_attention_output_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_10_attention_output_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(225541568)))]; tensor linear_63_cast_fp16 = linear(bias = text_model_encoder_layer_10_attention_output_dense_bias_to_fp16, weight = text_model_encoder_layer_10_attention_output_dense_weight_to_fp16, x = var_857_cast_fp16)[name = tensor("linear_63_cast_fp16")]; tensor input_215_cast_fp16 = add(x = linear_63_cast_fp16, y = hidden_states_59_cast_fp16)[name = tensor("input_215_cast_fp16")]; tensor input_217_axes_0 = const()[name = tensor("input_217_axes_0"), val = tensor([-1])]; tensor text_model_encoder_layer_10_attention_output_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_10_attention_output_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(225543168)))]; tensor text_model_encoder_layer_10_attention_output_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_10_attention_output_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(225544768)))]; tensor input_217_cast_fp16 = layer_norm(axes = input_217_axes_0, beta = text_model_encoder_layer_10_attention_output_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_encoder_layer_10_attention_output_LayerNorm_weight_to_fp16, x = input_215_cast_fp16)[name = tensor("input_217_cast_fp16")]; tensor text_model_encoder_layer_10_intermediate_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_10_intermediate_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(225546368)))]; tensor text_model_encoder_layer_10_intermediate_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_10_intermediate_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(230265024)))]; tensor linear_64_cast_fp16 = linear(bias = text_model_encoder_layer_10_intermediate_dense_bias_to_fp16, weight = text_model_encoder_layer_10_intermediate_dense_weight_to_fp16, x = input_217_cast_fp16)[name = tensor("linear_64_cast_fp16")]; tensor input_221_mode_0 = const()[name = tensor("input_221_mode_0"), val = tensor("EXACT")]; tensor input_221_cast_fp16 = gelu(mode = input_221_mode_0, x = linear_64_cast_fp16)[name = tensor("input_221_cast_fp16")]; tensor text_model_encoder_layer_10_output_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_10_output_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(230271232)))]; tensor text_model_encoder_layer_10_output_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_10_output_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(234989888)))]; tensor linear_65_cast_fp16 = linear(bias = text_model_encoder_layer_10_output_dense_bias_to_fp16, weight = text_model_encoder_layer_10_output_dense_weight_to_fp16, x = input_221_cast_fp16)[name = tensor("linear_65_cast_fp16")]; tensor input_225_cast_fp16 = add(x = linear_65_cast_fp16, y = input_217_cast_fp16)[name = tensor("input_225_cast_fp16")]; tensor hidden_states_65_axes_0 = const()[name = tensor("hidden_states_65_axes_0"), val = tensor([-1])]; tensor text_model_encoder_layer_10_output_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_10_output_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(234991488)))]; tensor text_model_encoder_layer_10_output_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_10_output_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(234993088)))]; tensor hidden_states_65_cast_fp16 = layer_norm(axes = hidden_states_65_axes_0, beta = text_model_encoder_layer_10_output_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_encoder_layer_10_output_LayerNorm_weight_to_fp16, x = input_225_cast_fp16)[name = tensor("hidden_states_65_cast_fp16")]; tensor text_model_encoder_layer_11_attention_self_query_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_11_attention_self_query_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(234994688)))]; tensor text_model_encoder_layer_11_attention_self_query_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_11_attention_self_query_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(236174400)))]; tensor linear_66_cast_fp16 = linear(bias = text_model_encoder_layer_11_attention_self_query_bias_to_fp16, weight = text_model_encoder_layer_11_attention_self_query_weight_to_fp16, x = hidden_states_65_cast_fp16)[name = tensor("linear_66_cast_fp16")]; tensor var_899 = const()[name = tensor("op_899"), val = tensor([1, 512, -1, 64])]; tensor var_900_cast_fp16 = reshape(shape = var_899, x = linear_66_cast_fp16)[name = tensor("op_900_cast_fp16")]; tensor text_model_encoder_layer_11_attention_self_key_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_11_attention_self_key_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(236176000)))]; tensor text_model_encoder_layer_11_attention_self_key_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_11_attention_self_key_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(237355712)))]; tensor linear_67_cast_fp16 = linear(bias = text_model_encoder_layer_11_attention_self_key_bias_to_fp16, weight = text_model_encoder_layer_11_attention_self_key_weight_to_fp16, x = hidden_states_65_cast_fp16)[name = tensor("linear_67_cast_fp16")]; tensor var_905 = const()[name = tensor("op_905"), val = tensor([1, 512, -1, 64])]; tensor var_906_cast_fp16 = reshape(shape = var_905, x = linear_67_cast_fp16)[name = tensor("op_906_cast_fp16")]; tensor text_model_encoder_layer_11_attention_self_value_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_11_attention_self_value_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(237357312)))]; tensor text_model_encoder_layer_11_attention_self_value_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_11_attention_self_value_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(238537024)))]; tensor linear_68_cast_fp16 = linear(bias = text_model_encoder_layer_11_attention_self_value_bias_to_fp16, weight = text_model_encoder_layer_11_attention_self_value_weight_to_fp16, x = hidden_states_65_cast_fp16)[name = tensor("linear_68_cast_fp16")]; tensor var_911 = const()[name = tensor("op_911"), val = tensor([1, 512, -1, 64])]; tensor var_912_cast_fp16 = reshape(shape = var_911, x = linear_68_cast_fp16)[name = tensor("op_912_cast_fp16")]; tensor value_23_perm_0 = const()[name = tensor("value_23_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_915_transpose_x_0 = const()[name = tensor("op_915_transpose_x_0"), val = tensor(false)]; tensor var_915_transpose_y_0 = const()[name = tensor("op_915_transpose_y_0"), val = tensor(false)]; tensor transpose_59_perm_0 = const()[name = tensor("transpose_59_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_60_perm_0 = const()[name = tensor("transpose_60_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_60 = transpose(perm = transpose_60_perm_0, x = var_906_cast_fp16)[name = tensor("transpose_62")]; tensor transpose_59 = transpose(perm = transpose_59_perm_0, x = var_900_cast_fp16)[name = tensor("transpose_63")]; tensor var_915_cast_fp16 = matmul(transpose_x = var_915_transpose_x_0, transpose_y = var_915_transpose_y_0, x = transpose_59, y = transpose_60)[name = tensor("op_915_cast_fp16")]; tensor var_916_to_fp16 = const()[name = tensor("op_916_to_fp16"), val = tensor(0x1p-3)]; tensor attn_weights_45_cast_fp16 = mul(x = var_915_cast_fp16, y = var_916_to_fp16)[name = tensor("attn_weights_45_cast_fp16")]; tensor input_227_cast_fp16 = add(x = attn_weights_45_cast_fp16, y = attention_mask_cast_fp16)[name = tensor("input_227_cast_fp16")]; tensor var_919_cast_fp16 = softmax(axis = var_21, x = input_227_cast_fp16)[name = tensor("op_919_cast_fp16")]; tensor attn_output_45_transpose_x_0 = const()[name = tensor("attn_output_45_transpose_x_0"), val = tensor(false)]; tensor attn_output_45_transpose_y_0 = const()[name = tensor("attn_output_45_transpose_y_0"), val = tensor(false)]; tensor value_23_cast_fp16 = transpose(perm = value_23_perm_0, x = var_912_cast_fp16)[name = tensor("transpose_64")]; tensor attn_output_45_cast_fp16 = matmul(transpose_x = attn_output_45_transpose_x_0, transpose_y = attn_output_45_transpose_y_0, x = var_919_cast_fp16, y = value_23_cast_fp16)[name = tensor("attn_output_45_cast_fp16")]; tensor var_923_perm_0 = const()[name = tensor("op_923_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_925 = const()[name = tensor("op_925"), val = tensor([1, 512, -1])]; tensor var_923_cast_fp16 = transpose(perm = var_923_perm_0, x = attn_output_45_cast_fp16)[name = tensor("transpose_61")]; tensor var_926_cast_fp16 = reshape(shape = var_925, x = var_923_cast_fp16)[name = tensor("op_926_cast_fp16")]; tensor text_model_encoder_layer_11_attention_output_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_11_attention_output_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(238538624)))]; tensor text_model_encoder_layer_11_attention_output_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_11_attention_output_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(239718336)))]; tensor linear_69_cast_fp16 = linear(bias = text_model_encoder_layer_11_attention_output_dense_bias_to_fp16, weight = text_model_encoder_layer_11_attention_output_dense_weight_to_fp16, x = var_926_cast_fp16)[name = tensor("linear_69_cast_fp16")]; tensor input_235_cast_fp16 = add(x = linear_69_cast_fp16, y = hidden_states_65_cast_fp16)[name = tensor("input_235_cast_fp16")]; tensor input_237_axes_0 = const()[name = tensor("input_237_axes_0"), val = tensor([-1])]; tensor text_model_encoder_layer_11_attention_output_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_11_attention_output_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(239719936)))]; tensor text_model_encoder_layer_11_attention_output_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_11_attention_output_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(239721536)))]; tensor input_237_cast_fp16 = layer_norm(axes = input_237_axes_0, beta = text_model_encoder_layer_11_attention_output_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_encoder_layer_11_attention_output_LayerNorm_weight_to_fp16, x = input_235_cast_fp16)[name = tensor("input_237_cast_fp16")]; tensor text_model_encoder_layer_11_intermediate_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_11_intermediate_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(239723136)))]; tensor text_model_encoder_layer_11_intermediate_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_11_intermediate_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(244441792)))]; tensor linear_70_cast_fp16 = linear(bias = text_model_encoder_layer_11_intermediate_dense_bias_to_fp16, weight = text_model_encoder_layer_11_intermediate_dense_weight_to_fp16, x = input_237_cast_fp16)[name = tensor("linear_70_cast_fp16")]; tensor input_241_mode_0 = const()[name = tensor("input_241_mode_0"), val = tensor("EXACT")]; tensor input_241_cast_fp16 = gelu(mode = input_241_mode_0, x = linear_70_cast_fp16)[name = tensor("input_241_cast_fp16")]; tensor text_model_encoder_layer_11_output_dense_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_11_output_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(244448000)))]; tensor text_model_encoder_layer_11_output_dense_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_11_output_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(249166656)))]; tensor linear_71_cast_fp16 = linear(bias = text_model_encoder_layer_11_output_dense_bias_to_fp16, weight = text_model_encoder_layer_11_output_dense_weight_to_fp16, x = input_241_cast_fp16)[name = tensor("linear_71_cast_fp16")]; tensor input_245_cast_fp16 = add(x = linear_71_cast_fp16, y = input_237_cast_fp16)[name = tensor("input_245_cast_fp16")]; tensor hidden_states_axes_0 = const()[name = tensor("hidden_states_axes_0"), val = tensor([-1])]; tensor text_model_encoder_layer_11_output_LayerNorm_weight_to_fp16 = const()[name = tensor("text_model_encoder_layer_11_output_LayerNorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(249168256)))]; tensor text_model_encoder_layer_11_output_LayerNorm_bias_to_fp16 = const()[name = tensor("text_model_encoder_layer_11_output_LayerNorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(249169856)))]; tensor hidden_states_cast_fp16 = layer_norm(axes = hidden_states_axes_0, beta = text_model_encoder_layer_11_output_LayerNorm_bias_to_fp16, epsilon = var_23_to_fp16, gamma = text_model_encoder_layer_11_output_LayerNorm_weight_to_fp16, x = input_245_cast_fp16)[name = tensor("hidden_states_cast_fp16")]; tensor input_247_begin_0 = const()[name = tensor("input_247_begin_0"), val = tensor([0, 0, 0])]; tensor input_247_end_0 = const()[name = tensor("input_247_end_0"), val = tensor([1, 1, 768])]; tensor input_247_end_mask_0 = const()[name = tensor("input_247_end_mask_0"), val = tensor([true, false, true])]; tensor input_247_squeeze_mask_0 = const()[name = tensor("input_247_squeeze_mask_0"), val = tensor([false, true, false])]; tensor input_247_cast_fp16 = slice_by_index(begin = input_247_begin_0, end = input_247_end_0, end_mask = input_247_end_mask_0, squeeze_mask = input_247_squeeze_mask_0, x = hidden_states_cast_fp16)[name = tensor("input_247_cast_fp16")]; tensor text_model_pooler_dense_weight_to_fp16 = const()[name = tensor("text_model_pooler_dense_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(249171456)))]; tensor text_model_pooler_dense_bias_to_fp16 = const()[name = tensor("text_model_pooler_dense_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(250351168)))]; tensor linear_72_cast_fp16 = linear(bias = text_model_pooler_dense_bias_to_fp16, weight = text_model_pooler_dense_weight_to_fp16, x = input_247_cast_fp16)[name = tensor("linear_72_cast_fp16")]; tensor input_251_cast_fp16 = tanh(x = linear_72_cast_fp16)[name = tensor("input_251_cast_fp16")]; tensor text_projection_linear1_weight_to_fp16 = const()[name = tensor("text_projection_linear1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(250352768)))]; tensor text_projection_linear1_bias_to_fp16 = const()[name = tensor("text_projection_linear1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(251139264)))]; tensor linear_73_cast_fp16 = linear(bias = text_projection_linear1_bias_to_fp16, weight = text_projection_linear1_weight_to_fp16, x = input_251_cast_fp16)[name = tensor("linear_73_cast_fp16")]; tensor input_cast_fp16 = relu(x = linear_73_cast_fp16)[name = tensor("input_cast_fp16")]; tensor text_projection_linear2_weight_to_fp16 = const()[name = tensor("text_projection_linear2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(251140352)))]; tensor text_projection_linear2_bias_to_fp16 = const()[name = tensor("text_projection_linear2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(251664704)))]; tensor linear_74_cast_fp16 = linear(bias = text_projection_linear2_bias_to_fp16, weight = text_projection_linear2_weight_to_fp16, x = input_cast_fp16)[name = tensor("linear_74_cast_fp16")]; tensor linear_74_cast_fp16_to_fp32_dtype_0 = const()[name = tensor("linear_74_cast_fp16_to_fp32_dtype_0"), val = tensor("fp32")]; tensor text_embeds = cast(dtype = linear_74_cast_fp16_to_fp32_dtype_0, x = linear_74_cast_fp16)[name = tensor("cast_56")]; } -> (text_embeds); }