program(1.0) [buildInfo = dict, tensor>({{"coremlc-component-MIL", "3405.2.1"}, {"coremlc-version", "3404.23.1"}, {"coremltools-component-torch", "2.6.0"}, {"coremltools-source-dialect", "TorchScript"}, {"coremltools-version", "8.3.0"}})] { func main(tensor input) { tensor var_11 = const()[name = tensor("op_11"), val = tensor(-1)]; tensor hidden_states_1_axes_0 = const()[name = tensor("hidden_states_1_axes_0"), val = tensor([-1])]; tensor input_to_fp16_dtype_0 = const()[name = tensor("input_to_fp16_dtype_0"), val = tensor("fp16")]; tensor encoder_layers_0_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_0_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(64)))]; tensor encoder_layers_0_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_0_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(2432)))]; tensor var_5_to_fp16 = const()[name = tensor("op_5_to_fp16"), val = tensor(0x1.1p-20)]; tensor input_to_fp16 = cast(dtype = input_to_fp16_dtype_0, x = input)[name = tensor("cast_136")]; tensor hidden_states_1_cast_fp16 = layer_norm(axes = hidden_states_1_axes_0, beta = encoder_layers_0_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_0_layer_norm1_weight_to_fp16, x = input_to_fp16)[name = tensor("hidden_states_1_cast_fp16")]; tensor encoder_layers_0_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_0_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(4800)))]; tensor encoder_layers_0_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_0_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(2659072)))]; tensor linear_0_cast_fp16 = linear(bias = encoder_layers_0_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_0_self_attn_q_proj_weight_to_fp16, x = hidden_states_1_cast_fp16)[name = tensor("linear_0_cast_fp16")]; tensor encoder_layers_0_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_0_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(2661440)))]; tensor encoder_layers_0_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_0_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(5315712)))]; tensor linear_1_cast_fp16 = linear(bias = encoder_layers_0_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_0_self_attn_k_proj_weight_to_fp16, x = hidden_states_1_cast_fp16)[name = tensor("linear_1_cast_fp16")]; tensor encoder_layers_0_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_0_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(5318080)))]; tensor encoder_layers_0_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_0_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(7972352)))]; tensor linear_2_cast_fp16 = linear(bias = encoder_layers_0_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_0_self_attn_v_proj_weight_to_fp16, x = hidden_states_1_cast_fp16)[name = tensor("linear_2_cast_fp16")]; tensor var_96 = const()[name = tensor("op_96"), val = tensor([1, 1024, 16, 72])]; tensor var_97_cast_fp16 = reshape(shape = var_96, x = linear_0_cast_fp16)[name = tensor("op_97_cast_fp16")]; tensor var_99 = const()[name = tensor("op_99"), val = tensor([1, 1024, 16, 72])]; tensor var_100_cast_fp16 = reshape(shape = var_99, x = linear_1_cast_fp16)[name = tensor("op_100_cast_fp16")]; tensor var_102 = const()[name = tensor("op_102"), val = tensor([1, 1024, 16, 72])]; tensor var_103_cast_fp16 = reshape(shape = var_102, x = linear_2_cast_fp16)[name = tensor("op_103_cast_fp16")]; tensor value_states_3_perm_0 = const()[name = tensor("value_states_3_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_106_transpose_x_0 = const()[name = tensor("op_106_transpose_x_0"), val = tensor(false)]; tensor var_106_transpose_y_0 = const()[name = tensor("op_106_transpose_y_0"), val = tensor(false)]; tensor transpose_81_perm_0 = const()[name = tensor("transpose_81_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_82_perm_0 = const()[name = tensor("transpose_82_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_82 = transpose(perm = transpose_82_perm_0, x = var_100_cast_fp16)[name = tensor("transpose_240")]; tensor transpose_81 = transpose(perm = transpose_81_perm_0, x = var_97_cast_fp16)[name = tensor("transpose_241")]; tensor var_106_cast_fp16 = matmul(transpose_x = var_106_transpose_x_0, transpose_y = var_106_transpose_y_0, x = transpose_81, y = transpose_82)[name = tensor("op_106_cast_fp16")]; tensor var_107_to_fp16 = const()[name = tensor("op_107_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_1_cast_fp16 = mul(x = var_106_cast_fp16, y = var_107_to_fp16)[name = tensor("attn_weights_1_cast_fp16")]; tensor var_109_cast_fp16 = softmax(axis = var_11, x = attn_weights_1_cast_fp16)[name = tensor("op_109_cast_fp16")]; tensor attn_output_1_transpose_x_0 = const()[name = tensor("attn_output_1_transpose_x_0"), val = tensor(false)]; tensor attn_output_1_transpose_y_0 = const()[name = tensor("attn_output_1_transpose_y_0"), val = tensor(false)]; tensor value_states_3_cast_fp16 = transpose(perm = value_states_3_perm_0, x = var_103_cast_fp16)[name = tensor("transpose_242")]; tensor attn_output_1_cast_fp16 = matmul(transpose_x = attn_output_1_transpose_x_0, transpose_y = attn_output_1_transpose_y_0, x = var_109_cast_fp16, y = value_states_3_cast_fp16)[name = tensor("attn_output_1_cast_fp16")]; tensor var_113_perm_0 = const()[name = tensor("op_113_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_115 = const()[name = tensor("op_115"), val = tensor([1, 1024, 1152])]; tensor var_113_cast_fp16 = transpose(perm = var_113_perm_0, x = attn_output_1_cast_fp16)[name = tensor("transpose_239")]; tensor input_3_cast_fp16 = reshape(shape = var_115, x = var_113_cast_fp16)[name = tensor("input_3_cast_fp16")]; tensor encoder_layers_0_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_0_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(7974720)))]; tensor encoder_layers_0_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_0_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(10628992)))]; tensor linear_3_cast_fp16 = linear(bias = encoder_layers_0_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_0_self_attn_out_proj_weight_to_fp16, x = input_3_cast_fp16)[name = tensor("linear_3_cast_fp16")]; tensor input_5_cast_fp16 = add(x = input_to_fp16, y = linear_3_cast_fp16)[name = tensor("input_5_cast_fp16")]; tensor input_7_axes_0 = const()[name = tensor("input_7_axes_0"), val = tensor([-1])]; tensor encoder_layers_0_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_0_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(10631360)))]; tensor encoder_layers_0_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_0_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(10633728)))]; tensor input_7_cast_fp16 = layer_norm(axes = input_7_axes_0, beta = encoder_layers_0_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_0_layer_norm2_weight_to_fp16, x = input_5_cast_fp16)[name = tensor("input_7_cast_fp16")]; tensor encoder_layers_0_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_0_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(10636096)))]; tensor encoder_layers_0_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_0_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(20552576)))]; tensor linear_4_cast_fp16 = linear(bias = encoder_layers_0_mlp_fc1_bias_to_fp16, weight = encoder_layers_0_mlp_fc1_weight_to_fp16, x = input_7_cast_fp16)[name = tensor("linear_4_cast_fp16")]; tensor input_11_mode_0 = const()[name = tensor("input_11_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_11_cast_fp16 = gelu(mode = input_11_mode_0, x = linear_4_cast_fp16)[name = tensor("input_11_cast_fp16")]; tensor encoder_layers_0_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_0_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(20561280)))]; tensor encoder_layers_0_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_0_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(30477760)))]; tensor linear_5_cast_fp16 = linear(bias = encoder_layers_0_mlp_fc2_bias_to_fp16, weight = encoder_layers_0_mlp_fc2_weight_to_fp16, x = input_11_cast_fp16)[name = tensor("linear_5_cast_fp16")]; tensor input_13_cast_fp16 = add(x = input_5_cast_fp16, y = linear_5_cast_fp16)[name = tensor("input_13_cast_fp16")]; tensor hidden_states_7_axes_0 = const()[name = tensor("hidden_states_7_axes_0"), val = tensor([-1])]; tensor encoder_layers_1_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_1_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(30480128)))]; tensor encoder_layers_1_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_1_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(30482496)))]; tensor hidden_states_7_cast_fp16 = layer_norm(axes = hidden_states_7_axes_0, beta = encoder_layers_1_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_1_layer_norm1_weight_to_fp16, x = input_13_cast_fp16)[name = tensor("hidden_states_7_cast_fp16")]; tensor encoder_layers_1_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_1_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(30484864)))]; tensor encoder_layers_1_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_1_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(33139136)))]; tensor linear_6_cast_fp16 = linear(bias = encoder_layers_1_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_1_self_attn_q_proj_weight_to_fp16, x = hidden_states_7_cast_fp16)[name = tensor("linear_6_cast_fp16")]; tensor encoder_layers_1_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_1_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(33141504)))]; tensor encoder_layers_1_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_1_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(35795776)))]; tensor linear_7_cast_fp16 = linear(bias = encoder_layers_1_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_1_self_attn_k_proj_weight_to_fp16, x = hidden_states_7_cast_fp16)[name = tensor("linear_7_cast_fp16")]; tensor encoder_layers_1_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_1_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(35798144)))]; tensor encoder_layers_1_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_1_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(38452416)))]; tensor linear_8_cast_fp16 = linear(bias = encoder_layers_1_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_1_self_attn_v_proj_weight_to_fp16, x = hidden_states_7_cast_fp16)[name = tensor("linear_8_cast_fp16")]; tensor var_158 = const()[name = tensor("op_158"), val = tensor([1, 1024, 16, 72])]; tensor var_159_cast_fp16 = reshape(shape = var_158, x = linear_6_cast_fp16)[name = tensor("op_159_cast_fp16")]; tensor var_161 = const()[name = tensor("op_161"), val = tensor([1, 1024, 16, 72])]; tensor var_162_cast_fp16 = reshape(shape = var_161, x = linear_7_cast_fp16)[name = tensor("op_162_cast_fp16")]; tensor var_164 = const()[name = tensor("op_164"), val = tensor([1, 1024, 16, 72])]; tensor var_165_cast_fp16 = reshape(shape = var_164, x = linear_8_cast_fp16)[name = tensor("op_165_cast_fp16")]; tensor value_states_7_perm_0 = const()[name = tensor("value_states_7_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_168_transpose_x_0 = const()[name = tensor("op_168_transpose_x_0"), val = tensor(false)]; tensor var_168_transpose_y_0 = const()[name = tensor("op_168_transpose_y_0"), val = tensor(false)]; tensor transpose_83_perm_0 = const()[name = tensor("transpose_83_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_84_perm_0 = const()[name = tensor("transpose_84_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_84 = transpose(perm = transpose_84_perm_0, x = var_162_cast_fp16)[name = tensor("transpose_236")]; tensor transpose_83 = transpose(perm = transpose_83_perm_0, x = var_159_cast_fp16)[name = tensor("transpose_237")]; tensor var_168_cast_fp16 = matmul(transpose_x = var_168_transpose_x_0, transpose_y = var_168_transpose_y_0, x = transpose_83, y = transpose_84)[name = tensor("op_168_cast_fp16")]; tensor var_169_to_fp16 = const()[name = tensor("op_169_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_5_cast_fp16 = mul(x = var_168_cast_fp16, y = var_169_to_fp16)[name = tensor("attn_weights_5_cast_fp16")]; tensor var_171_cast_fp16 = softmax(axis = var_11, x = attn_weights_5_cast_fp16)[name = tensor("op_171_cast_fp16")]; tensor attn_output_5_transpose_x_0 = const()[name = tensor("attn_output_5_transpose_x_0"), val = tensor(false)]; tensor attn_output_5_transpose_y_0 = const()[name = tensor("attn_output_5_transpose_y_0"), val = tensor(false)]; tensor value_states_7_cast_fp16 = transpose(perm = value_states_7_perm_0, x = var_165_cast_fp16)[name = tensor("transpose_238")]; tensor attn_output_5_cast_fp16 = matmul(transpose_x = attn_output_5_transpose_x_0, transpose_y = attn_output_5_transpose_y_0, x = var_171_cast_fp16, y = value_states_7_cast_fp16)[name = tensor("attn_output_5_cast_fp16")]; tensor var_175_perm_0 = const()[name = tensor("op_175_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_177 = const()[name = tensor("op_177"), val = tensor([1, 1024, 1152])]; tensor var_175_cast_fp16 = transpose(perm = var_175_perm_0, x = attn_output_5_cast_fp16)[name = tensor("transpose_235")]; tensor input_17_cast_fp16 = reshape(shape = var_177, x = var_175_cast_fp16)[name = tensor("input_17_cast_fp16")]; tensor encoder_layers_1_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_1_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(38454784)))]; tensor encoder_layers_1_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_1_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(41109056)))]; tensor linear_9_cast_fp16 = linear(bias = encoder_layers_1_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_1_self_attn_out_proj_weight_to_fp16, x = input_17_cast_fp16)[name = tensor("linear_9_cast_fp16")]; tensor input_19_cast_fp16 = add(x = input_13_cast_fp16, y = linear_9_cast_fp16)[name = tensor("input_19_cast_fp16")]; tensor input_21_axes_0 = const()[name = tensor("input_21_axes_0"), val = tensor([-1])]; tensor encoder_layers_1_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_1_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(41111424)))]; tensor encoder_layers_1_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_1_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(41113792)))]; tensor input_21_cast_fp16 = layer_norm(axes = input_21_axes_0, beta = encoder_layers_1_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_1_layer_norm2_weight_to_fp16, x = input_19_cast_fp16)[name = tensor("input_21_cast_fp16")]; tensor encoder_layers_1_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_1_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(41116160)))]; tensor encoder_layers_1_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_1_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(51032640)))]; tensor linear_10_cast_fp16 = linear(bias = encoder_layers_1_mlp_fc1_bias_to_fp16, weight = encoder_layers_1_mlp_fc1_weight_to_fp16, x = input_21_cast_fp16)[name = tensor("linear_10_cast_fp16")]; tensor input_25_mode_0 = const()[name = tensor("input_25_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_25_cast_fp16 = gelu(mode = input_25_mode_0, x = linear_10_cast_fp16)[name = tensor("input_25_cast_fp16")]; tensor encoder_layers_1_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_1_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(51041344)))]; tensor encoder_layers_1_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_1_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(60957824)))]; tensor linear_11_cast_fp16 = linear(bias = encoder_layers_1_mlp_fc2_bias_to_fp16, weight = encoder_layers_1_mlp_fc2_weight_to_fp16, x = input_25_cast_fp16)[name = tensor("linear_11_cast_fp16")]; tensor input_27_cast_fp16 = add(x = input_19_cast_fp16, y = linear_11_cast_fp16)[name = tensor("input_27_cast_fp16")]; tensor hidden_states_13_axes_0 = const()[name = tensor("hidden_states_13_axes_0"), val = tensor([-1])]; tensor encoder_layers_2_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_2_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(60960192)))]; tensor encoder_layers_2_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_2_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(60962560)))]; tensor hidden_states_13_cast_fp16 = layer_norm(axes = hidden_states_13_axes_0, beta = encoder_layers_2_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_2_layer_norm1_weight_to_fp16, x = input_27_cast_fp16)[name = tensor("hidden_states_13_cast_fp16")]; tensor encoder_layers_2_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_2_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(60964928)))]; tensor encoder_layers_2_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_2_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(63619200)))]; tensor linear_12_cast_fp16 = linear(bias = encoder_layers_2_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_2_self_attn_q_proj_weight_to_fp16, x = hidden_states_13_cast_fp16)[name = tensor("linear_12_cast_fp16")]; tensor encoder_layers_2_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_2_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(63621568)))]; tensor encoder_layers_2_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_2_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(66275840)))]; tensor linear_13_cast_fp16 = linear(bias = encoder_layers_2_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_2_self_attn_k_proj_weight_to_fp16, x = hidden_states_13_cast_fp16)[name = tensor("linear_13_cast_fp16")]; tensor encoder_layers_2_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_2_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(66278208)))]; tensor encoder_layers_2_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_2_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(68932480)))]; tensor linear_14_cast_fp16 = linear(bias = encoder_layers_2_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_2_self_attn_v_proj_weight_to_fp16, x = hidden_states_13_cast_fp16)[name = tensor("linear_14_cast_fp16")]; tensor var_220 = const()[name = tensor("op_220"), val = tensor([1, 1024, 16, 72])]; tensor var_221_cast_fp16 = reshape(shape = var_220, x = linear_12_cast_fp16)[name = tensor("op_221_cast_fp16")]; tensor var_223 = const()[name = tensor("op_223"), val = tensor([1, 1024, 16, 72])]; tensor var_224_cast_fp16 = reshape(shape = var_223, x = linear_13_cast_fp16)[name = tensor("op_224_cast_fp16")]; tensor var_226 = const()[name = tensor("op_226"), val = tensor([1, 1024, 16, 72])]; tensor var_227_cast_fp16 = reshape(shape = var_226, x = linear_14_cast_fp16)[name = tensor("op_227_cast_fp16")]; tensor value_states_11_perm_0 = const()[name = tensor("value_states_11_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_230_transpose_x_0 = const()[name = tensor("op_230_transpose_x_0"), val = tensor(false)]; tensor var_230_transpose_y_0 = const()[name = tensor("op_230_transpose_y_0"), val = tensor(false)]; tensor transpose_85_perm_0 = const()[name = tensor("transpose_85_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_86_perm_0 = const()[name = tensor("transpose_86_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_86 = transpose(perm = transpose_86_perm_0, x = var_224_cast_fp16)[name = tensor("transpose_232")]; tensor transpose_85 = transpose(perm = transpose_85_perm_0, x = var_221_cast_fp16)[name = tensor("transpose_233")]; tensor var_230_cast_fp16 = matmul(transpose_x = var_230_transpose_x_0, transpose_y = var_230_transpose_y_0, x = transpose_85, y = transpose_86)[name = tensor("op_230_cast_fp16")]; tensor var_231_to_fp16 = const()[name = tensor("op_231_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_9_cast_fp16 = mul(x = var_230_cast_fp16, y = var_231_to_fp16)[name = tensor("attn_weights_9_cast_fp16")]; tensor var_233_cast_fp16 = softmax(axis = var_11, x = attn_weights_9_cast_fp16)[name = tensor("op_233_cast_fp16")]; tensor attn_output_9_transpose_x_0 = const()[name = tensor("attn_output_9_transpose_x_0"), val = tensor(false)]; tensor attn_output_9_transpose_y_0 = const()[name = tensor("attn_output_9_transpose_y_0"), val = tensor(false)]; tensor value_states_11_cast_fp16 = transpose(perm = value_states_11_perm_0, x = var_227_cast_fp16)[name = tensor("transpose_234")]; tensor attn_output_9_cast_fp16 = matmul(transpose_x = attn_output_9_transpose_x_0, transpose_y = attn_output_9_transpose_y_0, x = var_233_cast_fp16, y = value_states_11_cast_fp16)[name = tensor("attn_output_9_cast_fp16")]; tensor var_237_perm_0 = const()[name = tensor("op_237_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_239 = const()[name = tensor("op_239"), val = tensor([1, 1024, 1152])]; tensor var_237_cast_fp16 = transpose(perm = var_237_perm_0, x = attn_output_9_cast_fp16)[name = tensor("transpose_231")]; tensor input_31_cast_fp16 = reshape(shape = var_239, x = var_237_cast_fp16)[name = tensor("input_31_cast_fp16")]; tensor encoder_layers_2_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_2_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(68934848)))]; tensor encoder_layers_2_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_2_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(71589120)))]; tensor linear_15_cast_fp16 = linear(bias = encoder_layers_2_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_2_self_attn_out_proj_weight_to_fp16, x = input_31_cast_fp16)[name = tensor("linear_15_cast_fp16")]; tensor input_33_cast_fp16 = add(x = input_27_cast_fp16, y = linear_15_cast_fp16)[name = tensor("input_33_cast_fp16")]; tensor input_35_axes_0 = const()[name = tensor("input_35_axes_0"), val = tensor([-1])]; tensor encoder_layers_2_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_2_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(71591488)))]; tensor encoder_layers_2_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_2_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(71593856)))]; tensor input_35_cast_fp16 = layer_norm(axes = input_35_axes_0, beta = encoder_layers_2_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_2_layer_norm2_weight_to_fp16, x = input_33_cast_fp16)[name = tensor("input_35_cast_fp16")]; tensor encoder_layers_2_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_2_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(71596224)))]; tensor encoder_layers_2_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_2_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(81512704)))]; tensor linear_16_cast_fp16 = linear(bias = encoder_layers_2_mlp_fc1_bias_to_fp16, weight = encoder_layers_2_mlp_fc1_weight_to_fp16, x = input_35_cast_fp16)[name = tensor("linear_16_cast_fp16")]; tensor input_39_mode_0 = const()[name = tensor("input_39_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_39_cast_fp16 = gelu(mode = input_39_mode_0, x = linear_16_cast_fp16)[name = tensor("input_39_cast_fp16")]; tensor encoder_layers_2_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_2_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(81521408)))]; tensor encoder_layers_2_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_2_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(91437888)))]; tensor linear_17_cast_fp16 = linear(bias = encoder_layers_2_mlp_fc2_bias_to_fp16, weight = encoder_layers_2_mlp_fc2_weight_to_fp16, x = input_39_cast_fp16)[name = tensor("linear_17_cast_fp16")]; tensor input_41_cast_fp16 = add(x = input_33_cast_fp16, y = linear_17_cast_fp16)[name = tensor("input_41_cast_fp16")]; tensor hidden_states_19_axes_0 = const()[name = tensor("hidden_states_19_axes_0"), val = tensor([-1])]; tensor encoder_layers_3_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_3_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(91440256)))]; tensor encoder_layers_3_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_3_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(91442624)))]; tensor hidden_states_19_cast_fp16 = layer_norm(axes = hidden_states_19_axes_0, beta = encoder_layers_3_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_3_layer_norm1_weight_to_fp16, x = input_41_cast_fp16)[name = tensor("hidden_states_19_cast_fp16")]; tensor encoder_layers_3_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_3_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(91444992)))]; tensor encoder_layers_3_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_3_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(94099264)))]; tensor linear_18_cast_fp16 = linear(bias = encoder_layers_3_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_3_self_attn_q_proj_weight_to_fp16, x = hidden_states_19_cast_fp16)[name = tensor("linear_18_cast_fp16")]; tensor encoder_layers_3_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_3_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(94101632)))]; tensor encoder_layers_3_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_3_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(96755904)))]; tensor linear_19_cast_fp16 = linear(bias = encoder_layers_3_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_3_self_attn_k_proj_weight_to_fp16, x = hidden_states_19_cast_fp16)[name = tensor("linear_19_cast_fp16")]; tensor encoder_layers_3_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_3_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(96758272)))]; tensor encoder_layers_3_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_3_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(99412544)))]; tensor linear_20_cast_fp16 = linear(bias = encoder_layers_3_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_3_self_attn_v_proj_weight_to_fp16, x = hidden_states_19_cast_fp16)[name = tensor("linear_20_cast_fp16")]; tensor var_282 = const()[name = tensor("op_282"), val = tensor([1, 1024, 16, 72])]; tensor var_283_cast_fp16 = reshape(shape = var_282, x = linear_18_cast_fp16)[name = tensor("op_283_cast_fp16")]; tensor var_285 = const()[name = tensor("op_285"), val = tensor([1, 1024, 16, 72])]; tensor var_286_cast_fp16 = reshape(shape = var_285, x = linear_19_cast_fp16)[name = tensor("op_286_cast_fp16")]; tensor var_288 = const()[name = tensor("op_288"), val = tensor([1, 1024, 16, 72])]; tensor var_289_cast_fp16 = reshape(shape = var_288, x = linear_20_cast_fp16)[name = tensor("op_289_cast_fp16")]; tensor value_states_15_perm_0 = const()[name = tensor("value_states_15_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_292_transpose_x_0 = const()[name = tensor("op_292_transpose_x_0"), val = tensor(false)]; tensor var_292_transpose_y_0 = const()[name = tensor("op_292_transpose_y_0"), val = tensor(false)]; tensor transpose_87_perm_0 = const()[name = tensor("transpose_87_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_88_perm_0 = const()[name = tensor("transpose_88_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_88 = transpose(perm = transpose_88_perm_0, x = var_286_cast_fp16)[name = tensor("transpose_228")]; tensor transpose_87 = transpose(perm = transpose_87_perm_0, x = var_283_cast_fp16)[name = tensor("transpose_229")]; tensor var_292_cast_fp16 = matmul(transpose_x = var_292_transpose_x_0, transpose_y = var_292_transpose_y_0, x = transpose_87, y = transpose_88)[name = tensor("op_292_cast_fp16")]; tensor var_293_to_fp16 = const()[name = tensor("op_293_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_13_cast_fp16 = mul(x = var_292_cast_fp16, y = var_293_to_fp16)[name = tensor("attn_weights_13_cast_fp16")]; tensor var_295_cast_fp16 = softmax(axis = var_11, x = attn_weights_13_cast_fp16)[name = tensor("op_295_cast_fp16")]; tensor attn_output_13_transpose_x_0 = const()[name = tensor("attn_output_13_transpose_x_0"), val = tensor(false)]; tensor attn_output_13_transpose_y_0 = const()[name = tensor("attn_output_13_transpose_y_0"), val = tensor(false)]; tensor value_states_15_cast_fp16 = transpose(perm = value_states_15_perm_0, x = var_289_cast_fp16)[name = tensor("transpose_230")]; tensor attn_output_13_cast_fp16 = matmul(transpose_x = attn_output_13_transpose_x_0, transpose_y = attn_output_13_transpose_y_0, x = var_295_cast_fp16, y = value_states_15_cast_fp16)[name = tensor("attn_output_13_cast_fp16")]; tensor var_299_perm_0 = const()[name = tensor("op_299_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_301 = const()[name = tensor("op_301"), val = tensor([1, 1024, 1152])]; tensor var_299_cast_fp16 = transpose(perm = var_299_perm_0, x = attn_output_13_cast_fp16)[name = tensor("transpose_227")]; tensor input_45_cast_fp16 = reshape(shape = var_301, x = var_299_cast_fp16)[name = tensor("input_45_cast_fp16")]; tensor encoder_layers_3_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_3_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(99414912)))]; tensor encoder_layers_3_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_3_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(102069184)))]; tensor linear_21_cast_fp16 = linear(bias = encoder_layers_3_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_3_self_attn_out_proj_weight_to_fp16, x = input_45_cast_fp16)[name = tensor("linear_21_cast_fp16")]; tensor input_47_cast_fp16 = add(x = input_41_cast_fp16, y = linear_21_cast_fp16)[name = tensor("input_47_cast_fp16")]; tensor input_49_axes_0 = const()[name = tensor("input_49_axes_0"), val = tensor([-1])]; tensor encoder_layers_3_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_3_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(102071552)))]; tensor encoder_layers_3_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_3_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(102073920)))]; tensor input_49_cast_fp16 = layer_norm(axes = input_49_axes_0, beta = encoder_layers_3_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_3_layer_norm2_weight_to_fp16, x = input_47_cast_fp16)[name = tensor("input_49_cast_fp16")]; tensor encoder_layers_3_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_3_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(102076288)))]; tensor encoder_layers_3_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_3_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(111992768)))]; tensor linear_22_cast_fp16 = linear(bias = encoder_layers_3_mlp_fc1_bias_to_fp16, weight = encoder_layers_3_mlp_fc1_weight_to_fp16, x = input_49_cast_fp16)[name = tensor("linear_22_cast_fp16")]; tensor input_53_mode_0 = const()[name = tensor("input_53_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_53_cast_fp16 = gelu(mode = input_53_mode_0, x = linear_22_cast_fp16)[name = tensor("input_53_cast_fp16")]; tensor encoder_layers_3_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_3_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(112001472)))]; tensor encoder_layers_3_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_3_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(121917952)))]; tensor linear_23_cast_fp16 = linear(bias = encoder_layers_3_mlp_fc2_bias_to_fp16, weight = encoder_layers_3_mlp_fc2_weight_to_fp16, x = input_53_cast_fp16)[name = tensor("linear_23_cast_fp16")]; tensor input_55_cast_fp16 = add(x = input_47_cast_fp16, y = linear_23_cast_fp16)[name = tensor("input_55_cast_fp16")]; tensor hidden_states_25_axes_0 = const()[name = tensor("hidden_states_25_axes_0"), val = tensor([-1])]; tensor encoder_layers_4_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_4_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(121920320)))]; tensor encoder_layers_4_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_4_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(121922688)))]; tensor hidden_states_25_cast_fp16 = layer_norm(axes = hidden_states_25_axes_0, beta = encoder_layers_4_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_4_layer_norm1_weight_to_fp16, x = input_55_cast_fp16)[name = tensor("hidden_states_25_cast_fp16")]; tensor encoder_layers_4_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_4_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(121925056)))]; tensor encoder_layers_4_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_4_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(124579328)))]; tensor linear_24_cast_fp16 = linear(bias = encoder_layers_4_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_4_self_attn_q_proj_weight_to_fp16, x = hidden_states_25_cast_fp16)[name = tensor("linear_24_cast_fp16")]; tensor encoder_layers_4_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_4_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(124581696)))]; tensor encoder_layers_4_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_4_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(127235968)))]; tensor linear_25_cast_fp16 = linear(bias = encoder_layers_4_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_4_self_attn_k_proj_weight_to_fp16, x = hidden_states_25_cast_fp16)[name = tensor("linear_25_cast_fp16")]; tensor encoder_layers_4_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_4_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(127238336)))]; tensor encoder_layers_4_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_4_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(129892608)))]; tensor linear_26_cast_fp16 = linear(bias = encoder_layers_4_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_4_self_attn_v_proj_weight_to_fp16, x = hidden_states_25_cast_fp16)[name = tensor("linear_26_cast_fp16")]; tensor var_344 = const()[name = tensor("op_344"), val = tensor([1, 1024, 16, 72])]; tensor var_345_cast_fp16 = reshape(shape = var_344, x = linear_24_cast_fp16)[name = tensor("op_345_cast_fp16")]; tensor var_347 = const()[name = tensor("op_347"), val = tensor([1, 1024, 16, 72])]; tensor var_348_cast_fp16 = reshape(shape = var_347, x = linear_25_cast_fp16)[name = tensor("op_348_cast_fp16")]; tensor var_350 = const()[name = tensor("op_350"), val = tensor([1, 1024, 16, 72])]; tensor var_351_cast_fp16 = reshape(shape = var_350, x = linear_26_cast_fp16)[name = tensor("op_351_cast_fp16")]; tensor value_states_19_perm_0 = const()[name = tensor("value_states_19_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_354_transpose_x_0 = const()[name = tensor("op_354_transpose_x_0"), val = tensor(false)]; tensor var_354_transpose_y_0 = const()[name = tensor("op_354_transpose_y_0"), val = tensor(false)]; tensor transpose_89_perm_0 = const()[name = tensor("transpose_89_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_90_perm_0 = const()[name = tensor("transpose_90_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_90 = transpose(perm = transpose_90_perm_0, x = var_348_cast_fp16)[name = tensor("transpose_224")]; tensor transpose_89 = transpose(perm = transpose_89_perm_0, x = var_345_cast_fp16)[name = tensor("transpose_225")]; tensor var_354_cast_fp16 = matmul(transpose_x = var_354_transpose_x_0, transpose_y = var_354_transpose_y_0, x = transpose_89, y = transpose_90)[name = tensor("op_354_cast_fp16")]; tensor var_355_to_fp16 = const()[name = tensor("op_355_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_17_cast_fp16 = mul(x = var_354_cast_fp16, y = var_355_to_fp16)[name = tensor("attn_weights_17_cast_fp16")]; tensor var_357_cast_fp16 = softmax(axis = var_11, x = attn_weights_17_cast_fp16)[name = tensor("op_357_cast_fp16")]; tensor attn_output_17_transpose_x_0 = const()[name = tensor("attn_output_17_transpose_x_0"), val = tensor(false)]; tensor attn_output_17_transpose_y_0 = const()[name = tensor("attn_output_17_transpose_y_0"), val = tensor(false)]; tensor value_states_19_cast_fp16 = transpose(perm = value_states_19_perm_0, x = var_351_cast_fp16)[name = tensor("transpose_226")]; tensor attn_output_17_cast_fp16 = matmul(transpose_x = attn_output_17_transpose_x_0, transpose_y = attn_output_17_transpose_y_0, x = var_357_cast_fp16, y = value_states_19_cast_fp16)[name = tensor("attn_output_17_cast_fp16")]; tensor var_361_perm_0 = const()[name = tensor("op_361_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_363 = const()[name = tensor("op_363"), val = tensor([1, 1024, 1152])]; tensor var_361_cast_fp16 = transpose(perm = var_361_perm_0, x = attn_output_17_cast_fp16)[name = tensor("transpose_223")]; tensor input_59_cast_fp16 = reshape(shape = var_363, x = var_361_cast_fp16)[name = tensor("input_59_cast_fp16")]; tensor encoder_layers_4_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_4_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(129894976)))]; tensor encoder_layers_4_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_4_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(132549248)))]; tensor linear_27_cast_fp16 = linear(bias = encoder_layers_4_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_4_self_attn_out_proj_weight_to_fp16, x = input_59_cast_fp16)[name = tensor("linear_27_cast_fp16")]; tensor input_61_cast_fp16 = add(x = input_55_cast_fp16, y = linear_27_cast_fp16)[name = tensor("input_61_cast_fp16")]; tensor input_63_axes_0 = const()[name = tensor("input_63_axes_0"), val = tensor([-1])]; tensor encoder_layers_4_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_4_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(132551616)))]; tensor encoder_layers_4_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_4_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(132553984)))]; tensor input_63_cast_fp16 = layer_norm(axes = input_63_axes_0, beta = encoder_layers_4_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_4_layer_norm2_weight_to_fp16, x = input_61_cast_fp16)[name = tensor("input_63_cast_fp16")]; tensor encoder_layers_4_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_4_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(132556352)))]; tensor encoder_layers_4_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_4_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(142472832)))]; tensor linear_28_cast_fp16 = linear(bias = encoder_layers_4_mlp_fc1_bias_to_fp16, weight = encoder_layers_4_mlp_fc1_weight_to_fp16, x = input_63_cast_fp16)[name = tensor("linear_28_cast_fp16")]; tensor input_67_mode_0 = const()[name = tensor("input_67_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_67_cast_fp16 = gelu(mode = input_67_mode_0, x = linear_28_cast_fp16)[name = tensor("input_67_cast_fp16")]; tensor encoder_layers_4_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_4_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(142481536)))]; tensor encoder_layers_4_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_4_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(152398016)))]; tensor linear_29_cast_fp16 = linear(bias = encoder_layers_4_mlp_fc2_bias_to_fp16, weight = encoder_layers_4_mlp_fc2_weight_to_fp16, x = input_67_cast_fp16)[name = tensor("linear_29_cast_fp16")]; tensor input_69_cast_fp16 = add(x = input_61_cast_fp16, y = linear_29_cast_fp16)[name = tensor("input_69_cast_fp16")]; tensor hidden_states_31_axes_0 = const()[name = tensor("hidden_states_31_axes_0"), val = tensor([-1])]; tensor encoder_layers_5_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_5_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(152400384)))]; tensor encoder_layers_5_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_5_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(152402752)))]; tensor hidden_states_31_cast_fp16 = layer_norm(axes = hidden_states_31_axes_0, beta = encoder_layers_5_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_5_layer_norm1_weight_to_fp16, x = input_69_cast_fp16)[name = tensor("hidden_states_31_cast_fp16")]; tensor encoder_layers_5_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_5_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(152405120)))]; tensor encoder_layers_5_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_5_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(155059392)))]; tensor linear_30_cast_fp16 = linear(bias = encoder_layers_5_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_5_self_attn_q_proj_weight_to_fp16, x = hidden_states_31_cast_fp16)[name = tensor("linear_30_cast_fp16")]; tensor encoder_layers_5_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_5_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(155061760)))]; tensor encoder_layers_5_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_5_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(157716032)))]; tensor linear_31_cast_fp16 = linear(bias = encoder_layers_5_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_5_self_attn_k_proj_weight_to_fp16, x = hidden_states_31_cast_fp16)[name = tensor("linear_31_cast_fp16")]; tensor encoder_layers_5_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_5_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(157718400)))]; tensor encoder_layers_5_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_5_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(160372672)))]; tensor linear_32_cast_fp16 = linear(bias = encoder_layers_5_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_5_self_attn_v_proj_weight_to_fp16, x = hidden_states_31_cast_fp16)[name = tensor("linear_32_cast_fp16")]; tensor var_406 = const()[name = tensor("op_406"), val = tensor([1, 1024, 16, 72])]; tensor var_407_cast_fp16 = reshape(shape = var_406, x = linear_30_cast_fp16)[name = tensor("op_407_cast_fp16")]; tensor var_409 = const()[name = tensor("op_409"), val = tensor([1, 1024, 16, 72])]; tensor var_410_cast_fp16 = reshape(shape = var_409, x = linear_31_cast_fp16)[name = tensor("op_410_cast_fp16")]; tensor var_412 = const()[name = tensor("op_412"), val = tensor([1, 1024, 16, 72])]; tensor var_413_cast_fp16 = reshape(shape = var_412, x = linear_32_cast_fp16)[name = tensor("op_413_cast_fp16")]; tensor value_states_23_perm_0 = const()[name = tensor("value_states_23_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_416_transpose_x_0 = const()[name = tensor("op_416_transpose_x_0"), val = tensor(false)]; tensor var_416_transpose_y_0 = const()[name = tensor("op_416_transpose_y_0"), val = tensor(false)]; tensor transpose_91_perm_0 = const()[name = tensor("transpose_91_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_92_perm_0 = const()[name = tensor("transpose_92_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_92 = transpose(perm = transpose_92_perm_0, x = var_410_cast_fp16)[name = tensor("transpose_220")]; tensor transpose_91 = transpose(perm = transpose_91_perm_0, x = var_407_cast_fp16)[name = tensor("transpose_221")]; tensor var_416_cast_fp16 = matmul(transpose_x = var_416_transpose_x_0, transpose_y = var_416_transpose_y_0, x = transpose_91, y = transpose_92)[name = tensor("op_416_cast_fp16")]; tensor var_417_to_fp16 = const()[name = tensor("op_417_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_21_cast_fp16 = mul(x = var_416_cast_fp16, y = var_417_to_fp16)[name = tensor("attn_weights_21_cast_fp16")]; tensor var_419_cast_fp16 = softmax(axis = var_11, x = attn_weights_21_cast_fp16)[name = tensor("op_419_cast_fp16")]; tensor attn_output_21_transpose_x_0 = const()[name = tensor("attn_output_21_transpose_x_0"), val = tensor(false)]; tensor attn_output_21_transpose_y_0 = const()[name = tensor("attn_output_21_transpose_y_0"), val = tensor(false)]; tensor value_states_23_cast_fp16 = transpose(perm = value_states_23_perm_0, x = var_413_cast_fp16)[name = tensor("transpose_222")]; tensor attn_output_21_cast_fp16 = matmul(transpose_x = attn_output_21_transpose_x_0, transpose_y = attn_output_21_transpose_y_0, x = var_419_cast_fp16, y = value_states_23_cast_fp16)[name = tensor("attn_output_21_cast_fp16")]; tensor var_423_perm_0 = const()[name = tensor("op_423_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_425 = const()[name = tensor("op_425"), val = tensor([1, 1024, 1152])]; tensor var_423_cast_fp16 = transpose(perm = var_423_perm_0, x = attn_output_21_cast_fp16)[name = tensor("transpose_219")]; tensor input_73_cast_fp16 = reshape(shape = var_425, x = var_423_cast_fp16)[name = tensor("input_73_cast_fp16")]; tensor encoder_layers_5_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_5_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(160375040)))]; tensor encoder_layers_5_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_5_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(163029312)))]; tensor linear_33_cast_fp16 = linear(bias = encoder_layers_5_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_5_self_attn_out_proj_weight_to_fp16, x = input_73_cast_fp16)[name = tensor("linear_33_cast_fp16")]; tensor input_75_cast_fp16 = add(x = input_69_cast_fp16, y = linear_33_cast_fp16)[name = tensor("input_75_cast_fp16")]; tensor input_77_axes_0 = const()[name = tensor("input_77_axes_0"), val = tensor([-1])]; tensor encoder_layers_5_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_5_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(163031680)))]; tensor encoder_layers_5_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_5_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(163034048)))]; tensor input_77_cast_fp16 = layer_norm(axes = input_77_axes_0, beta = encoder_layers_5_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_5_layer_norm2_weight_to_fp16, x = input_75_cast_fp16)[name = tensor("input_77_cast_fp16")]; tensor encoder_layers_5_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_5_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(163036416)))]; tensor encoder_layers_5_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_5_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(172952896)))]; tensor linear_34_cast_fp16 = linear(bias = encoder_layers_5_mlp_fc1_bias_to_fp16, weight = encoder_layers_5_mlp_fc1_weight_to_fp16, x = input_77_cast_fp16)[name = tensor("linear_34_cast_fp16")]; tensor input_81_mode_0 = const()[name = tensor("input_81_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_81_cast_fp16 = gelu(mode = input_81_mode_0, x = linear_34_cast_fp16)[name = tensor("input_81_cast_fp16")]; tensor encoder_layers_5_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_5_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(172961600)))]; tensor encoder_layers_5_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_5_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(182878080)))]; tensor linear_35_cast_fp16 = linear(bias = encoder_layers_5_mlp_fc2_bias_to_fp16, weight = encoder_layers_5_mlp_fc2_weight_to_fp16, x = input_81_cast_fp16)[name = tensor("linear_35_cast_fp16")]; tensor input_83_cast_fp16 = add(x = input_75_cast_fp16, y = linear_35_cast_fp16)[name = tensor("input_83_cast_fp16")]; tensor hidden_states_37_axes_0 = const()[name = tensor("hidden_states_37_axes_0"), val = tensor([-1])]; tensor encoder_layers_6_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_6_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(182880448)))]; tensor encoder_layers_6_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_6_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(182882816)))]; tensor hidden_states_37_cast_fp16 = layer_norm(axes = hidden_states_37_axes_0, beta = encoder_layers_6_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_6_layer_norm1_weight_to_fp16, x = input_83_cast_fp16)[name = tensor("hidden_states_37_cast_fp16")]; tensor encoder_layers_6_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_6_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(182885184)))]; tensor encoder_layers_6_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_6_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(185539456)))]; tensor linear_36_cast_fp16 = linear(bias = encoder_layers_6_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_6_self_attn_q_proj_weight_to_fp16, x = hidden_states_37_cast_fp16)[name = tensor("linear_36_cast_fp16")]; tensor encoder_layers_6_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_6_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(185541824)))]; tensor encoder_layers_6_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_6_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(188196096)))]; tensor linear_37_cast_fp16 = linear(bias = encoder_layers_6_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_6_self_attn_k_proj_weight_to_fp16, x = hidden_states_37_cast_fp16)[name = tensor("linear_37_cast_fp16")]; tensor encoder_layers_6_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_6_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(188198464)))]; tensor encoder_layers_6_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_6_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(190852736)))]; tensor linear_38_cast_fp16 = linear(bias = encoder_layers_6_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_6_self_attn_v_proj_weight_to_fp16, x = hidden_states_37_cast_fp16)[name = tensor("linear_38_cast_fp16")]; tensor var_468 = const()[name = tensor("op_468"), val = tensor([1, 1024, 16, 72])]; tensor var_469_cast_fp16 = reshape(shape = var_468, x = linear_36_cast_fp16)[name = tensor("op_469_cast_fp16")]; tensor var_471 = const()[name = tensor("op_471"), val = tensor([1, 1024, 16, 72])]; tensor var_472_cast_fp16 = reshape(shape = var_471, x = linear_37_cast_fp16)[name = tensor("op_472_cast_fp16")]; tensor var_474 = const()[name = tensor("op_474"), val = tensor([1, 1024, 16, 72])]; tensor var_475_cast_fp16 = reshape(shape = var_474, x = linear_38_cast_fp16)[name = tensor("op_475_cast_fp16")]; tensor value_states_27_perm_0 = const()[name = tensor("value_states_27_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_478_transpose_x_0 = const()[name = tensor("op_478_transpose_x_0"), val = tensor(false)]; tensor var_478_transpose_y_0 = const()[name = tensor("op_478_transpose_y_0"), val = tensor(false)]; tensor transpose_93_perm_0 = const()[name = tensor("transpose_93_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_94_perm_0 = const()[name = tensor("transpose_94_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_94 = transpose(perm = transpose_94_perm_0, x = var_472_cast_fp16)[name = tensor("transpose_216")]; tensor transpose_93 = transpose(perm = transpose_93_perm_0, x = var_469_cast_fp16)[name = tensor("transpose_217")]; tensor var_478_cast_fp16 = matmul(transpose_x = var_478_transpose_x_0, transpose_y = var_478_transpose_y_0, x = transpose_93, y = transpose_94)[name = tensor("op_478_cast_fp16")]; tensor var_479_to_fp16 = const()[name = tensor("op_479_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_25_cast_fp16 = mul(x = var_478_cast_fp16, y = var_479_to_fp16)[name = tensor("attn_weights_25_cast_fp16")]; tensor var_481_cast_fp16 = softmax(axis = var_11, x = attn_weights_25_cast_fp16)[name = tensor("op_481_cast_fp16")]; tensor attn_output_25_transpose_x_0 = const()[name = tensor("attn_output_25_transpose_x_0"), val = tensor(false)]; tensor attn_output_25_transpose_y_0 = const()[name = tensor("attn_output_25_transpose_y_0"), val = tensor(false)]; tensor value_states_27_cast_fp16 = transpose(perm = value_states_27_perm_0, x = var_475_cast_fp16)[name = tensor("transpose_218")]; tensor attn_output_25_cast_fp16 = matmul(transpose_x = attn_output_25_transpose_x_0, transpose_y = attn_output_25_transpose_y_0, x = var_481_cast_fp16, y = value_states_27_cast_fp16)[name = tensor("attn_output_25_cast_fp16")]; tensor var_485_perm_0 = const()[name = tensor("op_485_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_487 = const()[name = tensor("op_487"), val = tensor([1, 1024, 1152])]; tensor var_485_cast_fp16 = transpose(perm = var_485_perm_0, x = attn_output_25_cast_fp16)[name = tensor("transpose_215")]; tensor input_87_cast_fp16 = reshape(shape = var_487, x = var_485_cast_fp16)[name = tensor("input_87_cast_fp16")]; tensor encoder_layers_6_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_6_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(190855104)))]; tensor encoder_layers_6_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_6_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(193509376)))]; tensor linear_39_cast_fp16 = linear(bias = encoder_layers_6_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_6_self_attn_out_proj_weight_to_fp16, x = input_87_cast_fp16)[name = tensor("linear_39_cast_fp16")]; tensor input_89_cast_fp16 = add(x = input_83_cast_fp16, y = linear_39_cast_fp16)[name = tensor("input_89_cast_fp16")]; tensor input_91_axes_0 = const()[name = tensor("input_91_axes_0"), val = tensor([-1])]; tensor encoder_layers_6_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_6_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(193511744)))]; tensor encoder_layers_6_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_6_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(193514112)))]; tensor input_91_cast_fp16 = layer_norm(axes = input_91_axes_0, beta = encoder_layers_6_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_6_layer_norm2_weight_to_fp16, x = input_89_cast_fp16)[name = tensor("input_91_cast_fp16")]; tensor encoder_layers_6_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_6_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(193516480)))]; tensor encoder_layers_6_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_6_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(203432960)))]; tensor linear_40_cast_fp16 = linear(bias = encoder_layers_6_mlp_fc1_bias_to_fp16, weight = encoder_layers_6_mlp_fc1_weight_to_fp16, x = input_91_cast_fp16)[name = tensor("linear_40_cast_fp16")]; tensor input_95_mode_0 = const()[name = tensor("input_95_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_95_cast_fp16 = gelu(mode = input_95_mode_0, x = linear_40_cast_fp16)[name = tensor("input_95_cast_fp16")]; tensor encoder_layers_6_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_6_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(203441664)))]; tensor encoder_layers_6_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_6_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(213358144)))]; tensor linear_41_cast_fp16 = linear(bias = encoder_layers_6_mlp_fc2_bias_to_fp16, weight = encoder_layers_6_mlp_fc2_weight_to_fp16, x = input_95_cast_fp16)[name = tensor("linear_41_cast_fp16")]; tensor input_97_cast_fp16 = add(x = input_89_cast_fp16, y = linear_41_cast_fp16)[name = tensor("input_97_cast_fp16")]; tensor hidden_states_43_axes_0 = const()[name = tensor("hidden_states_43_axes_0"), val = tensor([-1])]; tensor encoder_layers_7_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_7_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(213360512)))]; tensor encoder_layers_7_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_7_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(213362880)))]; tensor hidden_states_43_cast_fp16 = layer_norm(axes = hidden_states_43_axes_0, beta = encoder_layers_7_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_7_layer_norm1_weight_to_fp16, x = input_97_cast_fp16)[name = tensor("hidden_states_43_cast_fp16")]; tensor encoder_layers_7_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_7_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(213365248)))]; tensor encoder_layers_7_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_7_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(216019520)))]; tensor linear_42_cast_fp16 = linear(bias = encoder_layers_7_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_7_self_attn_q_proj_weight_to_fp16, x = hidden_states_43_cast_fp16)[name = tensor("linear_42_cast_fp16")]; tensor encoder_layers_7_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_7_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(216021888)))]; tensor encoder_layers_7_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_7_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(218676160)))]; tensor linear_43_cast_fp16 = linear(bias = encoder_layers_7_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_7_self_attn_k_proj_weight_to_fp16, x = hidden_states_43_cast_fp16)[name = tensor("linear_43_cast_fp16")]; tensor encoder_layers_7_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_7_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(218678528)))]; tensor encoder_layers_7_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_7_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(221332800)))]; tensor linear_44_cast_fp16 = linear(bias = encoder_layers_7_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_7_self_attn_v_proj_weight_to_fp16, x = hidden_states_43_cast_fp16)[name = tensor("linear_44_cast_fp16")]; tensor var_530 = const()[name = tensor("op_530"), val = tensor([1, 1024, 16, 72])]; tensor var_531_cast_fp16 = reshape(shape = var_530, x = linear_42_cast_fp16)[name = tensor("op_531_cast_fp16")]; tensor var_533 = const()[name = tensor("op_533"), val = tensor([1, 1024, 16, 72])]; tensor var_534_cast_fp16 = reshape(shape = var_533, x = linear_43_cast_fp16)[name = tensor("op_534_cast_fp16")]; tensor var_536 = const()[name = tensor("op_536"), val = tensor([1, 1024, 16, 72])]; tensor var_537_cast_fp16 = reshape(shape = var_536, x = linear_44_cast_fp16)[name = tensor("op_537_cast_fp16")]; tensor value_states_31_perm_0 = const()[name = tensor("value_states_31_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_540_transpose_x_0 = const()[name = tensor("op_540_transpose_x_0"), val = tensor(false)]; tensor var_540_transpose_y_0 = const()[name = tensor("op_540_transpose_y_0"), val = tensor(false)]; tensor transpose_95_perm_0 = const()[name = tensor("transpose_95_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_96_perm_0 = const()[name = tensor("transpose_96_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_96 = transpose(perm = transpose_96_perm_0, x = var_534_cast_fp16)[name = tensor("transpose_212")]; tensor transpose_95 = transpose(perm = transpose_95_perm_0, x = var_531_cast_fp16)[name = tensor("transpose_213")]; tensor var_540_cast_fp16 = matmul(transpose_x = var_540_transpose_x_0, transpose_y = var_540_transpose_y_0, x = transpose_95, y = transpose_96)[name = tensor("op_540_cast_fp16")]; tensor var_541_to_fp16 = const()[name = tensor("op_541_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_29_cast_fp16 = mul(x = var_540_cast_fp16, y = var_541_to_fp16)[name = tensor("attn_weights_29_cast_fp16")]; tensor var_543_cast_fp16 = softmax(axis = var_11, x = attn_weights_29_cast_fp16)[name = tensor("op_543_cast_fp16")]; tensor attn_output_29_transpose_x_0 = const()[name = tensor("attn_output_29_transpose_x_0"), val = tensor(false)]; tensor attn_output_29_transpose_y_0 = const()[name = tensor("attn_output_29_transpose_y_0"), val = tensor(false)]; tensor value_states_31_cast_fp16 = transpose(perm = value_states_31_perm_0, x = var_537_cast_fp16)[name = tensor("transpose_214")]; tensor attn_output_29_cast_fp16 = matmul(transpose_x = attn_output_29_transpose_x_0, transpose_y = attn_output_29_transpose_y_0, x = var_543_cast_fp16, y = value_states_31_cast_fp16)[name = tensor("attn_output_29_cast_fp16")]; tensor var_547_perm_0 = const()[name = tensor("op_547_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_549 = const()[name = tensor("op_549"), val = tensor([1, 1024, 1152])]; tensor var_547_cast_fp16 = transpose(perm = var_547_perm_0, x = attn_output_29_cast_fp16)[name = tensor("transpose_211")]; tensor input_101_cast_fp16 = reshape(shape = var_549, x = var_547_cast_fp16)[name = tensor("input_101_cast_fp16")]; tensor encoder_layers_7_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_7_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(221335168)))]; tensor encoder_layers_7_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_7_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(223989440)))]; tensor linear_45_cast_fp16 = linear(bias = encoder_layers_7_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_7_self_attn_out_proj_weight_to_fp16, x = input_101_cast_fp16)[name = tensor("linear_45_cast_fp16")]; tensor input_103_cast_fp16 = add(x = input_97_cast_fp16, y = linear_45_cast_fp16)[name = tensor("input_103_cast_fp16")]; tensor input_105_axes_0 = const()[name = tensor("input_105_axes_0"), val = tensor([-1])]; tensor encoder_layers_7_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_7_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(223991808)))]; tensor encoder_layers_7_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_7_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(223994176)))]; tensor input_105_cast_fp16 = layer_norm(axes = input_105_axes_0, beta = encoder_layers_7_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_7_layer_norm2_weight_to_fp16, x = input_103_cast_fp16)[name = tensor("input_105_cast_fp16")]; tensor encoder_layers_7_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_7_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(223996544)))]; tensor encoder_layers_7_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_7_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(233913024)))]; tensor linear_46_cast_fp16 = linear(bias = encoder_layers_7_mlp_fc1_bias_to_fp16, weight = encoder_layers_7_mlp_fc1_weight_to_fp16, x = input_105_cast_fp16)[name = tensor("linear_46_cast_fp16")]; tensor input_109_mode_0 = const()[name = tensor("input_109_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_109_cast_fp16 = gelu(mode = input_109_mode_0, x = linear_46_cast_fp16)[name = tensor("input_109_cast_fp16")]; tensor encoder_layers_7_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_7_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(233921728)))]; tensor encoder_layers_7_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_7_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(243838208)))]; tensor linear_47_cast_fp16 = linear(bias = encoder_layers_7_mlp_fc2_bias_to_fp16, weight = encoder_layers_7_mlp_fc2_weight_to_fp16, x = input_109_cast_fp16)[name = tensor("linear_47_cast_fp16")]; tensor input_111_cast_fp16 = add(x = input_103_cast_fp16, y = linear_47_cast_fp16)[name = tensor("input_111_cast_fp16")]; tensor hidden_states_49_axes_0 = const()[name = tensor("hidden_states_49_axes_0"), val = tensor([-1])]; tensor encoder_layers_8_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_8_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(243840576)))]; tensor encoder_layers_8_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_8_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(243842944)))]; tensor hidden_states_49_cast_fp16 = layer_norm(axes = hidden_states_49_axes_0, beta = encoder_layers_8_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_8_layer_norm1_weight_to_fp16, x = input_111_cast_fp16)[name = tensor("hidden_states_49_cast_fp16")]; tensor encoder_layers_8_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_8_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(243845312)))]; tensor encoder_layers_8_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_8_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(246499584)))]; tensor linear_48_cast_fp16 = linear(bias = encoder_layers_8_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_8_self_attn_q_proj_weight_to_fp16, x = hidden_states_49_cast_fp16)[name = tensor("linear_48_cast_fp16")]; tensor encoder_layers_8_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_8_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(246501952)))]; tensor encoder_layers_8_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_8_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(249156224)))]; tensor linear_49_cast_fp16 = linear(bias = encoder_layers_8_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_8_self_attn_k_proj_weight_to_fp16, x = hidden_states_49_cast_fp16)[name = tensor("linear_49_cast_fp16")]; tensor encoder_layers_8_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_8_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(249158592)))]; tensor encoder_layers_8_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_8_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(251812864)))]; tensor linear_50_cast_fp16 = linear(bias = encoder_layers_8_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_8_self_attn_v_proj_weight_to_fp16, x = hidden_states_49_cast_fp16)[name = tensor("linear_50_cast_fp16")]; tensor var_592 = const()[name = tensor("op_592"), val = tensor([1, 1024, 16, 72])]; tensor var_593_cast_fp16 = reshape(shape = var_592, x = linear_48_cast_fp16)[name = tensor("op_593_cast_fp16")]; tensor var_595 = const()[name = tensor("op_595"), val = tensor([1, 1024, 16, 72])]; tensor var_596_cast_fp16 = reshape(shape = var_595, x = linear_49_cast_fp16)[name = tensor("op_596_cast_fp16")]; tensor var_598 = const()[name = tensor("op_598"), val = tensor([1, 1024, 16, 72])]; tensor var_599_cast_fp16 = reshape(shape = var_598, x = linear_50_cast_fp16)[name = tensor("op_599_cast_fp16")]; tensor value_states_35_perm_0 = const()[name = tensor("value_states_35_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_602_transpose_x_0 = const()[name = tensor("op_602_transpose_x_0"), val = tensor(false)]; tensor var_602_transpose_y_0 = const()[name = tensor("op_602_transpose_y_0"), val = tensor(false)]; tensor transpose_97_perm_0 = const()[name = tensor("transpose_97_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_98_perm_0 = const()[name = tensor("transpose_98_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_98 = transpose(perm = transpose_98_perm_0, x = var_596_cast_fp16)[name = tensor("transpose_208")]; tensor transpose_97 = transpose(perm = transpose_97_perm_0, x = var_593_cast_fp16)[name = tensor("transpose_209")]; tensor var_602_cast_fp16 = matmul(transpose_x = var_602_transpose_x_0, transpose_y = var_602_transpose_y_0, x = transpose_97, y = transpose_98)[name = tensor("op_602_cast_fp16")]; tensor var_603_to_fp16 = const()[name = tensor("op_603_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_33_cast_fp16 = mul(x = var_602_cast_fp16, y = var_603_to_fp16)[name = tensor("attn_weights_33_cast_fp16")]; tensor var_605_cast_fp16 = softmax(axis = var_11, x = attn_weights_33_cast_fp16)[name = tensor("op_605_cast_fp16")]; tensor attn_output_33_transpose_x_0 = const()[name = tensor("attn_output_33_transpose_x_0"), val = tensor(false)]; tensor attn_output_33_transpose_y_0 = const()[name = tensor("attn_output_33_transpose_y_0"), val = tensor(false)]; tensor value_states_35_cast_fp16 = transpose(perm = value_states_35_perm_0, x = var_599_cast_fp16)[name = tensor("transpose_210")]; tensor attn_output_33_cast_fp16 = matmul(transpose_x = attn_output_33_transpose_x_0, transpose_y = attn_output_33_transpose_y_0, x = var_605_cast_fp16, y = value_states_35_cast_fp16)[name = tensor("attn_output_33_cast_fp16")]; tensor var_609_perm_0 = const()[name = tensor("op_609_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_611 = const()[name = tensor("op_611"), val = tensor([1, 1024, 1152])]; tensor var_609_cast_fp16 = transpose(perm = var_609_perm_0, x = attn_output_33_cast_fp16)[name = tensor("transpose_207")]; tensor input_115_cast_fp16 = reshape(shape = var_611, x = var_609_cast_fp16)[name = tensor("input_115_cast_fp16")]; tensor encoder_layers_8_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_8_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(251815232)))]; tensor encoder_layers_8_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_8_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(254469504)))]; tensor linear_51_cast_fp16 = linear(bias = encoder_layers_8_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_8_self_attn_out_proj_weight_to_fp16, x = input_115_cast_fp16)[name = tensor("linear_51_cast_fp16")]; tensor input_117_cast_fp16 = add(x = input_111_cast_fp16, y = linear_51_cast_fp16)[name = tensor("input_117_cast_fp16")]; tensor input_119_axes_0 = const()[name = tensor("input_119_axes_0"), val = tensor([-1])]; tensor encoder_layers_8_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_8_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(254471872)))]; tensor encoder_layers_8_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_8_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(254474240)))]; tensor input_119_cast_fp16 = layer_norm(axes = input_119_axes_0, beta = encoder_layers_8_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_8_layer_norm2_weight_to_fp16, x = input_117_cast_fp16)[name = tensor("input_119_cast_fp16")]; tensor encoder_layers_8_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_8_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(254476608)))]; tensor encoder_layers_8_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_8_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(264393088)))]; tensor linear_52_cast_fp16 = linear(bias = encoder_layers_8_mlp_fc1_bias_to_fp16, weight = encoder_layers_8_mlp_fc1_weight_to_fp16, x = input_119_cast_fp16)[name = tensor("linear_52_cast_fp16")]; tensor input_123_mode_0 = const()[name = tensor("input_123_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_123_cast_fp16 = gelu(mode = input_123_mode_0, x = linear_52_cast_fp16)[name = tensor("input_123_cast_fp16")]; tensor encoder_layers_8_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_8_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(264401792)))]; tensor encoder_layers_8_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_8_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(274318272)))]; tensor linear_53_cast_fp16 = linear(bias = encoder_layers_8_mlp_fc2_bias_to_fp16, weight = encoder_layers_8_mlp_fc2_weight_to_fp16, x = input_123_cast_fp16)[name = tensor("linear_53_cast_fp16")]; tensor input_125_cast_fp16 = add(x = input_117_cast_fp16, y = linear_53_cast_fp16)[name = tensor("input_125_cast_fp16")]; tensor hidden_states_55_axes_0 = const()[name = tensor("hidden_states_55_axes_0"), val = tensor([-1])]; tensor encoder_layers_9_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_9_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(274320640)))]; tensor encoder_layers_9_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_9_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(274323008)))]; tensor hidden_states_55_cast_fp16 = layer_norm(axes = hidden_states_55_axes_0, beta = encoder_layers_9_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_9_layer_norm1_weight_to_fp16, x = input_125_cast_fp16)[name = tensor("hidden_states_55_cast_fp16")]; tensor encoder_layers_9_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_9_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(274325376)))]; tensor encoder_layers_9_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_9_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(276979648)))]; tensor linear_54_cast_fp16 = linear(bias = encoder_layers_9_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_9_self_attn_q_proj_weight_to_fp16, x = hidden_states_55_cast_fp16)[name = tensor("linear_54_cast_fp16")]; tensor encoder_layers_9_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_9_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(276982016)))]; tensor encoder_layers_9_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_9_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(279636288)))]; tensor linear_55_cast_fp16 = linear(bias = encoder_layers_9_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_9_self_attn_k_proj_weight_to_fp16, x = hidden_states_55_cast_fp16)[name = tensor("linear_55_cast_fp16")]; tensor encoder_layers_9_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_9_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(279638656)))]; tensor encoder_layers_9_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_9_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(282292928)))]; tensor linear_56_cast_fp16 = linear(bias = encoder_layers_9_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_9_self_attn_v_proj_weight_to_fp16, x = hidden_states_55_cast_fp16)[name = tensor("linear_56_cast_fp16")]; tensor var_654 = const()[name = tensor("op_654"), val = tensor([1, 1024, 16, 72])]; tensor var_655_cast_fp16 = reshape(shape = var_654, x = linear_54_cast_fp16)[name = tensor("op_655_cast_fp16")]; tensor var_657 = const()[name = tensor("op_657"), val = tensor([1, 1024, 16, 72])]; tensor var_658_cast_fp16 = reshape(shape = var_657, x = linear_55_cast_fp16)[name = tensor("op_658_cast_fp16")]; tensor var_660 = const()[name = tensor("op_660"), val = tensor([1, 1024, 16, 72])]; tensor var_661_cast_fp16 = reshape(shape = var_660, x = linear_56_cast_fp16)[name = tensor("op_661_cast_fp16")]; tensor value_states_39_perm_0 = const()[name = tensor("value_states_39_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_664_transpose_x_0 = const()[name = tensor("op_664_transpose_x_0"), val = tensor(false)]; tensor var_664_transpose_y_0 = const()[name = tensor("op_664_transpose_y_0"), val = tensor(false)]; tensor transpose_99_perm_0 = const()[name = tensor("transpose_99_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_100_perm_0 = const()[name = tensor("transpose_100_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_100 = transpose(perm = transpose_100_perm_0, x = var_658_cast_fp16)[name = tensor("transpose_204")]; tensor transpose_99 = transpose(perm = transpose_99_perm_0, x = var_655_cast_fp16)[name = tensor("transpose_205")]; tensor var_664_cast_fp16 = matmul(transpose_x = var_664_transpose_x_0, transpose_y = var_664_transpose_y_0, x = transpose_99, y = transpose_100)[name = tensor("op_664_cast_fp16")]; tensor var_665_to_fp16 = const()[name = tensor("op_665_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_37_cast_fp16 = mul(x = var_664_cast_fp16, y = var_665_to_fp16)[name = tensor("attn_weights_37_cast_fp16")]; tensor var_667_cast_fp16 = softmax(axis = var_11, x = attn_weights_37_cast_fp16)[name = tensor("op_667_cast_fp16")]; tensor attn_output_37_transpose_x_0 = const()[name = tensor("attn_output_37_transpose_x_0"), val = tensor(false)]; tensor attn_output_37_transpose_y_0 = const()[name = tensor("attn_output_37_transpose_y_0"), val = tensor(false)]; tensor value_states_39_cast_fp16 = transpose(perm = value_states_39_perm_0, x = var_661_cast_fp16)[name = tensor("transpose_206")]; tensor attn_output_37_cast_fp16 = matmul(transpose_x = attn_output_37_transpose_x_0, transpose_y = attn_output_37_transpose_y_0, x = var_667_cast_fp16, y = value_states_39_cast_fp16)[name = tensor("attn_output_37_cast_fp16")]; tensor var_671_perm_0 = const()[name = tensor("op_671_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_673 = const()[name = tensor("op_673"), val = tensor([1, 1024, 1152])]; tensor var_671_cast_fp16 = transpose(perm = var_671_perm_0, x = attn_output_37_cast_fp16)[name = tensor("transpose_203")]; tensor input_129_cast_fp16 = reshape(shape = var_673, x = var_671_cast_fp16)[name = tensor("input_129_cast_fp16")]; tensor encoder_layers_9_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_9_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(282295296)))]; tensor encoder_layers_9_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_9_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(284949568)))]; tensor linear_57_cast_fp16 = linear(bias = encoder_layers_9_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_9_self_attn_out_proj_weight_to_fp16, x = input_129_cast_fp16)[name = tensor("linear_57_cast_fp16")]; tensor input_131_cast_fp16 = add(x = input_125_cast_fp16, y = linear_57_cast_fp16)[name = tensor("input_131_cast_fp16")]; tensor input_133_axes_0 = const()[name = tensor("input_133_axes_0"), val = tensor([-1])]; tensor encoder_layers_9_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_9_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(284951936)))]; tensor encoder_layers_9_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_9_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(284954304)))]; tensor input_133_cast_fp16 = layer_norm(axes = input_133_axes_0, beta = encoder_layers_9_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_9_layer_norm2_weight_to_fp16, x = input_131_cast_fp16)[name = tensor("input_133_cast_fp16")]; tensor encoder_layers_9_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_9_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(284956672)))]; tensor encoder_layers_9_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_9_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(294873152)))]; tensor linear_58_cast_fp16 = linear(bias = encoder_layers_9_mlp_fc1_bias_to_fp16, weight = encoder_layers_9_mlp_fc1_weight_to_fp16, x = input_133_cast_fp16)[name = tensor("linear_58_cast_fp16")]; tensor input_137_mode_0 = const()[name = tensor("input_137_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_137_cast_fp16 = gelu(mode = input_137_mode_0, x = linear_58_cast_fp16)[name = tensor("input_137_cast_fp16")]; tensor encoder_layers_9_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_9_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(294881856)))]; tensor encoder_layers_9_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_9_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(304798336)))]; tensor linear_59_cast_fp16 = linear(bias = encoder_layers_9_mlp_fc2_bias_to_fp16, weight = encoder_layers_9_mlp_fc2_weight_to_fp16, x = input_137_cast_fp16)[name = tensor("linear_59_cast_fp16")]; tensor input_139_cast_fp16 = add(x = input_131_cast_fp16, y = linear_59_cast_fp16)[name = tensor("input_139_cast_fp16")]; tensor hidden_states_61_axes_0 = const()[name = tensor("hidden_states_61_axes_0"), val = tensor([-1])]; tensor encoder_layers_10_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_10_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(304800704)))]; tensor encoder_layers_10_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_10_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(304803072)))]; tensor hidden_states_61_cast_fp16 = layer_norm(axes = hidden_states_61_axes_0, beta = encoder_layers_10_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_10_layer_norm1_weight_to_fp16, x = input_139_cast_fp16)[name = tensor("hidden_states_61_cast_fp16")]; tensor encoder_layers_10_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_10_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(304805440)))]; tensor encoder_layers_10_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_10_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(307459712)))]; tensor linear_60_cast_fp16 = linear(bias = encoder_layers_10_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_10_self_attn_q_proj_weight_to_fp16, x = hidden_states_61_cast_fp16)[name = tensor("linear_60_cast_fp16")]; tensor encoder_layers_10_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_10_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(307462080)))]; tensor encoder_layers_10_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_10_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(310116352)))]; tensor linear_61_cast_fp16 = linear(bias = encoder_layers_10_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_10_self_attn_k_proj_weight_to_fp16, x = hidden_states_61_cast_fp16)[name = tensor("linear_61_cast_fp16")]; tensor encoder_layers_10_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_10_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(310118720)))]; tensor encoder_layers_10_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_10_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(312772992)))]; tensor linear_62_cast_fp16 = linear(bias = encoder_layers_10_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_10_self_attn_v_proj_weight_to_fp16, x = hidden_states_61_cast_fp16)[name = tensor("linear_62_cast_fp16")]; tensor var_716 = const()[name = tensor("op_716"), val = tensor([1, 1024, 16, 72])]; tensor var_717_cast_fp16 = reshape(shape = var_716, x = linear_60_cast_fp16)[name = tensor("op_717_cast_fp16")]; tensor var_719 = const()[name = tensor("op_719"), val = tensor([1, 1024, 16, 72])]; tensor var_720_cast_fp16 = reshape(shape = var_719, x = linear_61_cast_fp16)[name = tensor("op_720_cast_fp16")]; tensor var_722 = const()[name = tensor("op_722"), val = tensor([1, 1024, 16, 72])]; tensor var_723_cast_fp16 = reshape(shape = var_722, x = linear_62_cast_fp16)[name = tensor("op_723_cast_fp16")]; tensor value_states_43_perm_0 = const()[name = tensor("value_states_43_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_726_transpose_x_0 = const()[name = tensor("op_726_transpose_x_0"), val = tensor(false)]; tensor var_726_transpose_y_0 = const()[name = tensor("op_726_transpose_y_0"), val = tensor(false)]; tensor transpose_101_perm_0 = const()[name = tensor("transpose_101_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_102_perm_0 = const()[name = tensor("transpose_102_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_102 = transpose(perm = transpose_102_perm_0, x = var_720_cast_fp16)[name = tensor("transpose_200")]; tensor transpose_101 = transpose(perm = transpose_101_perm_0, x = var_717_cast_fp16)[name = tensor("transpose_201")]; tensor var_726_cast_fp16 = matmul(transpose_x = var_726_transpose_x_0, transpose_y = var_726_transpose_y_0, x = transpose_101, y = transpose_102)[name = tensor("op_726_cast_fp16")]; tensor var_727_to_fp16 = const()[name = tensor("op_727_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_41_cast_fp16 = mul(x = var_726_cast_fp16, y = var_727_to_fp16)[name = tensor("attn_weights_41_cast_fp16")]; tensor var_729_cast_fp16 = softmax(axis = var_11, x = attn_weights_41_cast_fp16)[name = tensor("op_729_cast_fp16")]; tensor attn_output_41_transpose_x_0 = const()[name = tensor("attn_output_41_transpose_x_0"), val = tensor(false)]; tensor attn_output_41_transpose_y_0 = const()[name = tensor("attn_output_41_transpose_y_0"), val = tensor(false)]; tensor value_states_43_cast_fp16 = transpose(perm = value_states_43_perm_0, x = var_723_cast_fp16)[name = tensor("transpose_202")]; tensor attn_output_41_cast_fp16 = matmul(transpose_x = attn_output_41_transpose_x_0, transpose_y = attn_output_41_transpose_y_0, x = var_729_cast_fp16, y = value_states_43_cast_fp16)[name = tensor("attn_output_41_cast_fp16")]; tensor var_733_perm_0 = const()[name = tensor("op_733_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_735 = const()[name = tensor("op_735"), val = tensor([1, 1024, 1152])]; tensor var_733_cast_fp16 = transpose(perm = var_733_perm_0, x = attn_output_41_cast_fp16)[name = tensor("transpose_199")]; tensor input_143_cast_fp16 = reshape(shape = var_735, x = var_733_cast_fp16)[name = tensor("input_143_cast_fp16")]; tensor encoder_layers_10_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_10_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(312775360)))]; tensor encoder_layers_10_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_10_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(315429632)))]; tensor linear_63_cast_fp16 = linear(bias = encoder_layers_10_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_10_self_attn_out_proj_weight_to_fp16, x = input_143_cast_fp16)[name = tensor("linear_63_cast_fp16")]; tensor input_145_cast_fp16 = add(x = input_139_cast_fp16, y = linear_63_cast_fp16)[name = tensor("input_145_cast_fp16")]; tensor input_147_axes_0 = const()[name = tensor("input_147_axes_0"), val = tensor([-1])]; tensor encoder_layers_10_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_10_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(315432000)))]; tensor encoder_layers_10_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_10_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(315434368)))]; tensor input_147_cast_fp16 = layer_norm(axes = input_147_axes_0, beta = encoder_layers_10_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_10_layer_norm2_weight_to_fp16, x = input_145_cast_fp16)[name = tensor("input_147_cast_fp16")]; tensor encoder_layers_10_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_10_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(315436736)))]; tensor encoder_layers_10_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_10_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(325353216)))]; tensor linear_64_cast_fp16 = linear(bias = encoder_layers_10_mlp_fc1_bias_to_fp16, weight = encoder_layers_10_mlp_fc1_weight_to_fp16, x = input_147_cast_fp16)[name = tensor("linear_64_cast_fp16")]; tensor input_151_mode_0 = const()[name = tensor("input_151_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_151_cast_fp16 = gelu(mode = input_151_mode_0, x = linear_64_cast_fp16)[name = tensor("input_151_cast_fp16")]; tensor encoder_layers_10_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_10_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(325361920)))]; tensor encoder_layers_10_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_10_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(335278400)))]; tensor linear_65_cast_fp16 = linear(bias = encoder_layers_10_mlp_fc2_bias_to_fp16, weight = encoder_layers_10_mlp_fc2_weight_to_fp16, x = input_151_cast_fp16)[name = tensor("linear_65_cast_fp16")]; tensor input_153_cast_fp16 = add(x = input_145_cast_fp16, y = linear_65_cast_fp16)[name = tensor("input_153_cast_fp16")]; tensor hidden_states_67_axes_0 = const()[name = tensor("hidden_states_67_axes_0"), val = tensor([-1])]; tensor encoder_layers_11_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_11_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(335280768)))]; tensor encoder_layers_11_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_11_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(335283136)))]; tensor hidden_states_67_cast_fp16 = layer_norm(axes = hidden_states_67_axes_0, beta = encoder_layers_11_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_11_layer_norm1_weight_to_fp16, x = input_153_cast_fp16)[name = tensor("hidden_states_67_cast_fp16")]; tensor encoder_layers_11_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_11_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(335285504)))]; tensor encoder_layers_11_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_11_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(337939776)))]; tensor linear_66_cast_fp16 = linear(bias = encoder_layers_11_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_11_self_attn_q_proj_weight_to_fp16, x = hidden_states_67_cast_fp16)[name = tensor("linear_66_cast_fp16")]; tensor encoder_layers_11_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_11_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(337942144)))]; tensor encoder_layers_11_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_11_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(340596416)))]; tensor linear_67_cast_fp16 = linear(bias = encoder_layers_11_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_11_self_attn_k_proj_weight_to_fp16, x = hidden_states_67_cast_fp16)[name = tensor("linear_67_cast_fp16")]; tensor encoder_layers_11_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_11_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(340598784)))]; tensor encoder_layers_11_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_11_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(343253056)))]; tensor linear_68_cast_fp16 = linear(bias = encoder_layers_11_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_11_self_attn_v_proj_weight_to_fp16, x = hidden_states_67_cast_fp16)[name = tensor("linear_68_cast_fp16")]; tensor var_778 = const()[name = tensor("op_778"), val = tensor([1, 1024, 16, 72])]; tensor var_779_cast_fp16 = reshape(shape = var_778, x = linear_66_cast_fp16)[name = tensor("op_779_cast_fp16")]; tensor var_781 = const()[name = tensor("op_781"), val = tensor([1, 1024, 16, 72])]; tensor var_782_cast_fp16 = reshape(shape = var_781, x = linear_67_cast_fp16)[name = tensor("op_782_cast_fp16")]; tensor var_784 = const()[name = tensor("op_784"), val = tensor([1, 1024, 16, 72])]; tensor var_785_cast_fp16 = reshape(shape = var_784, x = linear_68_cast_fp16)[name = tensor("op_785_cast_fp16")]; tensor value_states_47_perm_0 = const()[name = tensor("value_states_47_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_788_transpose_x_0 = const()[name = tensor("op_788_transpose_x_0"), val = tensor(false)]; tensor var_788_transpose_y_0 = const()[name = tensor("op_788_transpose_y_0"), val = tensor(false)]; tensor transpose_103_perm_0 = const()[name = tensor("transpose_103_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_104_perm_0 = const()[name = tensor("transpose_104_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_104 = transpose(perm = transpose_104_perm_0, x = var_782_cast_fp16)[name = tensor("transpose_196")]; tensor transpose_103 = transpose(perm = transpose_103_perm_0, x = var_779_cast_fp16)[name = tensor("transpose_197")]; tensor var_788_cast_fp16 = matmul(transpose_x = var_788_transpose_x_0, transpose_y = var_788_transpose_y_0, x = transpose_103, y = transpose_104)[name = tensor("op_788_cast_fp16")]; tensor var_789_to_fp16 = const()[name = tensor("op_789_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_45_cast_fp16 = mul(x = var_788_cast_fp16, y = var_789_to_fp16)[name = tensor("attn_weights_45_cast_fp16")]; tensor var_791_cast_fp16 = softmax(axis = var_11, x = attn_weights_45_cast_fp16)[name = tensor("op_791_cast_fp16")]; tensor attn_output_45_transpose_x_0 = const()[name = tensor("attn_output_45_transpose_x_0"), val = tensor(false)]; tensor attn_output_45_transpose_y_0 = const()[name = tensor("attn_output_45_transpose_y_0"), val = tensor(false)]; tensor value_states_47_cast_fp16 = transpose(perm = value_states_47_perm_0, x = var_785_cast_fp16)[name = tensor("transpose_198")]; tensor attn_output_45_cast_fp16 = matmul(transpose_x = attn_output_45_transpose_x_0, transpose_y = attn_output_45_transpose_y_0, x = var_791_cast_fp16, y = value_states_47_cast_fp16)[name = tensor("attn_output_45_cast_fp16")]; tensor var_795_perm_0 = const()[name = tensor("op_795_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_797 = const()[name = tensor("op_797"), val = tensor([1, 1024, 1152])]; tensor var_795_cast_fp16 = transpose(perm = var_795_perm_0, x = attn_output_45_cast_fp16)[name = tensor("transpose_195")]; tensor input_157_cast_fp16 = reshape(shape = var_797, x = var_795_cast_fp16)[name = tensor("input_157_cast_fp16")]; tensor encoder_layers_11_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_11_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(343255424)))]; tensor encoder_layers_11_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_11_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(345909696)))]; tensor linear_69_cast_fp16 = linear(bias = encoder_layers_11_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_11_self_attn_out_proj_weight_to_fp16, x = input_157_cast_fp16)[name = tensor("linear_69_cast_fp16")]; tensor input_159_cast_fp16 = add(x = input_153_cast_fp16, y = linear_69_cast_fp16)[name = tensor("input_159_cast_fp16")]; tensor input_161_axes_0 = const()[name = tensor("input_161_axes_0"), val = tensor([-1])]; tensor encoder_layers_11_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_11_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(345912064)))]; tensor encoder_layers_11_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_11_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(345914432)))]; tensor input_161_cast_fp16 = layer_norm(axes = input_161_axes_0, beta = encoder_layers_11_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_11_layer_norm2_weight_to_fp16, x = input_159_cast_fp16)[name = tensor("input_161_cast_fp16")]; tensor encoder_layers_11_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_11_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(345916800)))]; tensor encoder_layers_11_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_11_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(355833280)))]; tensor linear_70_cast_fp16 = linear(bias = encoder_layers_11_mlp_fc1_bias_to_fp16, weight = encoder_layers_11_mlp_fc1_weight_to_fp16, x = input_161_cast_fp16)[name = tensor("linear_70_cast_fp16")]; tensor input_165_mode_0 = const()[name = tensor("input_165_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_165_cast_fp16 = gelu(mode = input_165_mode_0, x = linear_70_cast_fp16)[name = tensor("input_165_cast_fp16")]; tensor encoder_layers_11_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_11_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(355841984)))]; tensor encoder_layers_11_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_11_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(365758464)))]; tensor linear_71_cast_fp16 = linear(bias = encoder_layers_11_mlp_fc2_bias_to_fp16, weight = encoder_layers_11_mlp_fc2_weight_to_fp16, x = input_165_cast_fp16)[name = tensor("linear_71_cast_fp16")]; tensor input_167_cast_fp16 = add(x = input_159_cast_fp16, y = linear_71_cast_fp16)[name = tensor("input_167_cast_fp16")]; tensor hidden_states_73_axes_0 = const()[name = tensor("hidden_states_73_axes_0"), val = tensor([-1])]; tensor encoder_layers_12_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_12_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(365760832)))]; tensor encoder_layers_12_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_12_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(365763200)))]; tensor hidden_states_73_cast_fp16 = layer_norm(axes = hidden_states_73_axes_0, beta = encoder_layers_12_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_12_layer_norm1_weight_to_fp16, x = input_167_cast_fp16)[name = tensor("hidden_states_73_cast_fp16")]; tensor encoder_layers_12_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_12_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(365765568)))]; tensor encoder_layers_12_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_12_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(368419840)))]; tensor linear_72_cast_fp16 = linear(bias = encoder_layers_12_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_12_self_attn_q_proj_weight_to_fp16, x = hidden_states_73_cast_fp16)[name = tensor("linear_72_cast_fp16")]; tensor encoder_layers_12_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_12_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(368422208)))]; tensor encoder_layers_12_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_12_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(371076480)))]; tensor linear_73_cast_fp16 = linear(bias = encoder_layers_12_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_12_self_attn_k_proj_weight_to_fp16, x = hidden_states_73_cast_fp16)[name = tensor("linear_73_cast_fp16")]; tensor encoder_layers_12_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_12_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(371078848)))]; tensor encoder_layers_12_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_12_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(373733120)))]; tensor linear_74_cast_fp16 = linear(bias = encoder_layers_12_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_12_self_attn_v_proj_weight_to_fp16, x = hidden_states_73_cast_fp16)[name = tensor("linear_74_cast_fp16")]; tensor var_840 = const()[name = tensor("op_840"), val = tensor([1, 1024, 16, 72])]; tensor var_841_cast_fp16 = reshape(shape = var_840, x = linear_72_cast_fp16)[name = tensor("op_841_cast_fp16")]; tensor var_843 = const()[name = tensor("op_843"), val = tensor([1, 1024, 16, 72])]; tensor var_844_cast_fp16 = reshape(shape = var_843, x = linear_73_cast_fp16)[name = tensor("op_844_cast_fp16")]; tensor var_846 = const()[name = tensor("op_846"), val = tensor([1, 1024, 16, 72])]; tensor var_847_cast_fp16 = reshape(shape = var_846, x = linear_74_cast_fp16)[name = tensor("op_847_cast_fp16")]; tensor value_states_51_perm_0 = const()[name = tensor("value_states_51_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_850_transpose_x_0 = const()[name = tensor("op_850_transpose_x_0"), val = tensor(false)]; tensor var_850_transpose_y_0 = const()[name = tensor("op_850_transpose_y_0"), val = tensor(false)]; tensor transpose_105_perm_0 = const()[name = tensor("transpose_105_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_106_perm_0 = const()[name = tensor("transpose_106_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_106 = transpose(perm = transpose_106_perm_0, x = var_844_cast_fp16)[name = tensor("transpose_192")]; tensor transpose_105 = transpose(perm = transpose_105_perm_0, x = var_841_cast_fp16)[name = tensor("transpose_193")]; tensor var_850_cast_fp16 = matmul(transpose_x = var_850_transpose_x_0, transpose_y = var_850_transpose_y_0, x = transpose_105, y = transpose_106)[name = tensor("op_850_cast_fp16")]; tensor var_851_to_fp16 = const()[name = tensor("op_851_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_49_cast_fp16 = mul(x = var_850_cast_fp16, y = var_851_to_fp16)[name = tensor("attn_weights_49_cast_fp16")]; tensor var_853_cast_fp16 = softmax(axis = var_11, x = attn_weights_49_cast_fp16)[name = tensor("op_853_cast_fp16")]; tensor attn_output_49_transpose_x_0 = const()[name = tensor("attn_output_49_transpose_x_0"), val = tensor(false)]; tensor attn_output_49_transpose_y_0 = const()[name = tensor("attn_output_49_transpose_y_0"), val = tensor(false)]; tensor value_states_51_cast_fp16 = transpose(perm = value_states_51_perm_0, x = var_847_cast_fp16)[name = tensor("transpose_194")]; tensor attn_output_49_cast_fp16 = matmul(transpose_x = attn_output_49_transpose_x_0, transpose_y = attn_output_49_transpose_y_0, x = var_853_cast_fp16, y = value_states_51_cast_fp16)[name = tensor("attn_output_49_cast_fp16")]; tensor var_857_perm_0 = const()[name = tensor("op_857_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_859 = const()[name = tensor("op_859"), val = tensor([1, 1024, 1152])]; tensor var_857_cast_fp16 = transpose(perm = var_857_perm_0, x = attn_output_49_cast_fp16)[name = tensor("transpose_191")]; tensor input_171_cast_fp16 = reshape(shape = var_859, x = var_857_cast_fp16)[name = tensor("input_171_cast_fp16")]; tensor encoder_layers_12_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_12_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(373735488)))]; tensor encoder_layers_12_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_12_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(376389760)))]; tensor linear_75_cast_fp16 = linear(bias = encoder_layers_12_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_12_self_attn_out_proj_weight_to_fp16, x = input_171_cast_fp16)[name = tensor("linear_75_cast_fp16")]; tensor input_173_cast_fp16 = add(x = input_167_cast_fp16, y = linear_75_cast_fp16)[name = tensor("input_173_cast_fp16")]; tensor input_175_axes_0 = const()[name = tensor("input_175_axes_0"), val = tensor([-1])]; tensor encoder_layers_12_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_12_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(376392128)))]; tensor encoder_layers_12_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_12_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(376394496)))]; tensor input_175_cast_fp16 = layer_norm(axes = input_175_axes_0, beta = encoder_layers_12_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_12_layer_norm2_weight_to_fp16, x = input_173_cast_fp16)[name = tensor("input_175_cast_fp16")]; tensor encoder_layers_12_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_12_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(376396864)))]; tensor encoder_layers_12_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_12_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(386313344)))]; tensor linear_76_cast_fp16 = linear(bias = encoder_layers_12_mlp_fc1_bias_to_fp16, weight = encoder_layers_12_mlp_fc1_weight_to_fp16, x = input_175_cast_fp16)[name = tensor("linear_76_cast_fp16")]; tensor input_179_mode_0 = const()[name = tensor("input_179_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_179_cast_fp16 = gelu(mode = input_179_mode_0, x = linear_76_cast_fp16)[name = tensor("input_179_cast_fp16")]; tensor encoder_layers_12_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_12_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(386322048)))]; tensor encoder_layers_12_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_12_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(396238528)))]; tensor linear_77_cast_fp16 = linear(bias = encoder_layers_12_mlp_fc2_bias_to_fp16, weight = encoder_layers_12_mlp_fc2_weight_to_fp16, x = input_179_cast_fp16)[name = tensor("linear_77_cast_fp16")]; tensor input_181_cast_fp16 = add(x = input_173_cast_fp16, y = linear_77_cast_fp16)[name = tensor("input_181_cast_fp16")]; tensor hidden_states_79_axes_0 = const()[name = tensor("hidden_states_79_axes_0"), val = tensor([-1])]; tensor encoder_layers_13_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_13_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(396240896)))]; tensor encoder_layers_13_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_13_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(396243264)))]; tensor hidden_states_79_cast_fp16 = layer_norm(axes = hidden_states_79_axes_0, beta = encoder_layers_13_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_13_layer_norm1_weight_to_fp16, x = input_181_cast_fp16)[name = tensor("hidden_states_79_cast_fp16")]; tensor encoder_layers_13_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_13_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(396245632)))]; tensor encoder_layers_13_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_13_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(398899904)))]; tensor linear_78_cast_fp16 = linear(bias = encoder_layers_13_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_13_self_attn_q_proj_weight_to_fp16, x = hidden_states_79_cast_fp16)[name = tensor("linear_78_cast_fp16")]; tensor encoder_layers_13_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_13_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(398902272)))]; tensor encoder_layers_13_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_13_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(401556544)))]; tensor linear_79_cast_fp16 = linear(bias = encoder_layers_13_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_13_self_attn_k_proj_weight_to_fp16, x = hidden_states_79_cast_fp16)[name = tensor("linear_79_cast_fp16")]; tensor encoder_layers_13_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_13_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(401558912)))]; tensor encoder_layers_13_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_13_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(404213184)))]; tensor linear_80_cast_fp16 = linear(bias = encoder_layers_13_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_13_self_attn_v_proj_weight_to_fp16, x = hidden_states_79_cast_fp16)[name = tensor("linear_80_cast_fp16")]; tensor var_902 = const()[name = tensor("op_902"), val = tensor([1, 1024, 16, 72])]; tensor var_903_cast_fp16 = reshape(shape = var_902, x = linear_78_cast_fp16)[name = tensor("op_903_cast_fp16")]; tensor var_905 = const()[name = tensor("op_905"), val = tensor([1, 1024, 16, 72])]; tensor var_906_cast_fp16 = reshape(shape = var_905, x = linear_79_cast_fp16)[name = tensor("op_906_cast_fp16")]; tensor var_908 = const()[name = tensor("op_908"), val = tensor([1, 1024, 16, 72])]; tensor var_909_cast_fp16 = reshape(shape = var_908, x = linear_80_cast_fp16)[name = tensor("op_909_cast_fp16")]; tensor value_states_55_perm_0 = const()[name = tensor("value_states_55_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_912_transpose_x_0 = const()[name = tensor("op_912_transpose_x_0"), val = tensor(false)]; tensor var_912_transpose_y_0 = const()[name = tensor("op_912_transpose_y_0"), val = tensor(false)]; tensor transpose_107_perm_0 = const()[name = tensor("transpose_107_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_108_perm_0 = const()[name = tensor("transpose_108_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_108 = transpose(perm = transpose_108_perm_0, x = var_906_cast_fp16)[name = tensor("transpose_188")]; tensor transpose_107 = transpose(perm = transpose_107_perm_0, x = var_903_cast_fp16)[name = tensor("transpose_189")]; tensor var_912_cast_fp16 = matmul(transpose_x = var_912_transpose_x_0, transpose_y = var_912_transpose_y_0, x = transpose_107, y = transpose_108)[name = tensor("op_912_cast_fp16")]; tensor var_913_to_fp16 = const()[name = tensor("op_913_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_53_cast_fp16 = mul(x = var_912_cast_fp16, y = var_913_to_fp16)[name = tensor("attn_weights_53_cast_fp16")]; tensor var_915_cast_fp16 = softmax(axis = var_11, x = attn_weights_53_cast_fp16)[name = tensor("op_915_cast_fp16")]; tensor attn_output_53_transpose_x_0 = const()[name = tensor("attn_output_53_transpose_x_0"), val = tensor(false)]; tensor attn_output_53_transpose_y_0 = const()[name = tensor("attn_output_53_transpose_y_0"), val = tensor(false)]; tensor value_states_55_cast_fp16 = transpose(perm = value_states_55_perm_0, x = var_909_cast_fp16)[name = tensor("transpose_190")]; tensor attn_output_53_cast_fp16 = matmul(transpose_x = attn_output_53_transpose_x_0, transpose_y = attn_output_53_transpose_y_0, x = var_915_cast_fp16, y = value_states_55_cast_fp16)[name = tensor("attn_output_53_cast_fp16")]; tensor var_919_perm_0 = const()[name = tensor("op_919_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_921 = const()[name = tensor("op_921"), val = tensor([1, 1024, 1152])]; tensor var_919_cast_fp16 = transpose(perm = var_919_perm_0, x = attn_output_53_cast_fp16)[name = tensor("transpose_187")]; tensor input_185_cast_fp16 = reshape(shape = var_921, x = var_919_cast_fp16)[name = tensor("input_185_cast_fp16")]; tensor encoder_layers_13_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_13_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(404215552)))]; tensor encoder_layers_13_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_13_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(406869824)))]; tensor linear_81_cast_fp16 = linear(bias = encoder_layers_13_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_13_self_attn_out_proj_weight_to_fp16, x = input_185_cast_fp16)[name = tensor("linear_81_cast_fp16")]; tensor input_187_cast_fp16 = add(x = input_181_cast_fp16, y = linear_81_cast_fp16)[name = tensor("input_187_cast_fp16")]; tensor input_189_axes_0 = const()[name = tensor("input_189_axes_0"), val = tensor([-1])]; tensor encoder_layers_13_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_13_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(406872192)))]; tensor encoder_layers_13_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_13_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(406874560)))]; tensor input_189_cast_fp16 = layer_norm(axes = input_189_axes_0, beta = encoder_layers_13_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_13_layer_norm2_weight_to_fp16, x = input_187_cast_fp16)[name = tensor("input_189_cast_fp16")]; tensor encoder_layers_13_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_13_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(406876928)))]; tensor encoder_layers_13_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_13_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(416793408)))]; tensor linear_82_cast_fp16 = linear(bias = encoder_layers_13_mlp_fc1_bias_to_fp16, weight = encoder_layers_13_mlp_fc1_weight_to_fp16, x = input_189_cast_fp16)[name = tensor("linear_82_cast_fp16")]; tensor input_193_mode_0 = const()[name = tensor("input_193_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_193_cast_fp16 = gelu(mode = input_193_mode_0, x = linear_82_cast_fp16)[name = tensor("input_193_cast_fp16")]; tensor encoder_layers_13_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_13_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(416802112)))]; tensor encoder_layers_13_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_13_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(426718592)))]; tensor linear_83_cast_fp16 = linear(bias = encoder_layers_13_mlp_fc2_bias_to_fp16, weight = encoder_layers_13_mlp_fc2_weight_to_fp16, x = input_193_cast_fp16)[name = tensor("linear_83_cast_fp16")]; tensor input_195_cast_fp16 = add(x = input_187_cast_fp16, y = linear_83_cast_fp16)[name = tensor("input_195_cast_fp16")]; tensor hidden_states_85_axes_0 = const()[name = tensor("hidden_states_85_axes_0"), val = tensor([-1])]; tensor encoder_layers_14_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_14_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(426720960)))]; tensor encoder_layers_14_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_14_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(426723328)))]; tensor hidden_states_85_cast_fp16 = layer_norm(axes = hidden_states_85_axes_0, beta = encoder_layers_14_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_14_layer_norm1_weight_to_fp16, x = input_195_cast_fp16)[name = tensor("hidden_states_85_cast_fp16")]; tensor encoder_layers_14_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_14_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(426725696)))]; tensor encoder_layers_14_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_14_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(429379968)))]; tensor linear_84_cast_fp16 = linear(bias = encoder_layers_14_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_14_self_attn_q_proj_weight_to_fp16, x = hidden_states_85_cast_fp16)[name = tensor("linear_84_cast_fp16")]; tensor encoder_layers_14_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_14_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(429382336)))]; tensor encoder_layers_14_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_14_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(432036608)))]; tensor linear_85_cast_fp16 = linear(bias = encoder_layers_14_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_14_self_attn_k_proj_weight_to_fp16, x = hidden_states_85_cast_fp16)[name = tensor("linear_85_cast_fp16")]; tensor encoder_layers_14_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_14_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(432038976)))]; tensor encoder_layers_14_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_14_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(434693248)))]; tensor linear_86_cast_fp16 = linear(bias = encoder_layers_14_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_14_self_attn_v_proj_weight_to_fp16, x = hidden_states_85_cast_fp16)[name = tensor("linear_86_cast_fp16")]; tensor var_964 = const()[name = tensor("op_964"), val = tensor([1, 1024, 16, 72])]; tensor var_965_cast_fp16 = reshape(shape = var_964, x = linear_84_cast_fp16)[name = tensor("op_965_cast_fp16")]; tensor var_967 = const()[name = tensor("op_967"), val = tensor([1, 1024, 16, 72])]; tensor var_968_cast_fp16 = reshape(shape = var_967, x = linear_85_cast_fp16)[name = tensor("op_968_cast_fp16")]; tensor var_970 = const()[name = tensor("op_970"), val = tensor([1, 1024, 16, 72])]; tensor var_971_cast_fp16 = reshape(shape = var_970, x = linear_86_cast_fp16)[name = tensor("op_971_cast_fp16")]; tensor value_states_59_perm_0 = const()[name = tensor("value_states_59_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_974_transpose_x_0 = const()[name = tensor("op_974_transpose_x_0"), val = tensor(false)]; tensor var_974_transpose_y_0 = const()[name = tensor("op_974_transpose_y_0"), val = tensor(false)]; tensor transpose_109_perm_0 = const()[name = tensor("transpose_109_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_110_perm_0 = const()[name = tensor("transpose_110_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_110 = transpose(perm = transpose_110_perm_0, x = var_968_cast_fp16)[name = tensor("transpose_184")]; tensor transpose_109 = transpose(perm = transpose_109_perm_0, x = var_965_cast_fp16)[name = tensor("transpose_185")]; tensor var_974_cast_fp16 = matmul(transpose_x = var_974_transpose_x_0, transpose_y = var_974_transpose_y_0, x = transpose_109, y = transpose_110)[name = tensor("op_974_cast_fp16")]; tensor var_975_to_fp16 = const()[name = tensor("op_975_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_57_cast_fp16 = mul(x = var_974_cast_fp16, y = var_975_to_fp16)[name = tensor("attn_weights_57_cast_fp16")]; tensor var_977_cast_fp16 = softmax(axis = var_11, x = attn_weights_57_cast_fp16)[name = tensor("op_977_cast_fp16")]; tensor attn_output_57_transpose_x_0 = const()[name = tensor("attn_output_57_transpose_x_0"), val = tensor(false)]; tensor attn_output_57_transpose_y_0 = const()[name = tensor("attn_output_57_transpose_y_0"), val = tensor(false)]; tensor value_states_59_cast_fp16 = transpose(perm = value_states_59_perm_0, x = var_971_cast_fp16)[name = tensor("transpose_186")]; tensor attn_output_57_cast_fp16 = matmul(transpose_x = attn_output_57_transpose_x_0, transpose_y = attn_output_57_transpose_y_0, x = var_977_cast_fp16, y = value_states_59_cast_fp16)[name = tensor("attn_output_57_cast_fp16")]; tensor var_981_perm_0 = const()[name = tensor("op_981_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_983 = const()[name = tensor("op_983"), val = tensor([1, 1024, 1152])]; tensor var_981_cast_fp16 = transpose(perm = var_981_perm_0, x = attn_output_57_cast_fp16)[name = tensor("transpose_183")]; tensor input_199_cast_fp16 = reshape(shape = var_983, x = var_981_cast_fp16)[name = tensor("input_199_cast_fp16")]; tensor encoder_layers_14_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_14_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(434695616)))]; tensor encoder_layers_14_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_14_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(437349888)))]; tensor linear_87_cast_fp16 = linear(bias = encoder_layers_14_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_14_self_attn_out_proj_weight_to_fp16, x = input_199_cast_fp16)[name = tensor("linear_87_cast_fp16")]; tensor input_201_cast_fp16 = add(x = input_195_cast_fp16, y = linear_87_cast_fp16)[name = tensor("input_201_cast_fp16")]; tensor input_203_axes_0 = const()[name = tensor("input_203_axes_0"), val = tensor([-1])]; tensor encoder_layers_14_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_14_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(437352256)))]; tensor encoder_layers_14_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_14_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(437354624)))]; tensor input_203_cast_fp16 = layer_norm(axes = input_203_axes_0, beta = encoder_layers_14_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_14_layer_norm2_weight_to_fp16, x = input_201_cast_fp16)[name = tensor("input_203_cast_fp16")]; tensor encoder_layers_14_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_14_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(437356992)))]; tensor encoder_layers_14_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_14_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(447273472)))]; tensor linear_88_cast_fp16 = linear(bias = encoder_layers_14_mlp_fc1_bias_to_fp16, weight = encoder_layers_14_mlp_fc1_weight_to_fp16, x = input_203_cast_fp16)[name = tensor("linear_88_cast_fp16")]; tensor input_207_mode_0 = const()[name = tensor("input_207_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_207_cast_fp16 = gelu(mode = input_207_mode_0, x = linear_88_cast_fp16)[name = tensor("input_207_cast_fp16")]; tensor encoder_layers_14_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_14_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(447282176)))]; tensor encoder_layers_14_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_14_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(457198656)))]; tensor linear_89_cast_fp16 = linear(bias = encoder_layers_14_mlp_fc2_bias_to_fp16, weight = encoder_layers_14_mlp_fc2_weight_to_fp16, x = input_207_cast_fp16)[name = tensor("linear_89_cast_fp16")]; tensor input_209_cast_fp16 = add(x = input_201_cast_fp16, y = linear_89_cast_fp16)[name = tensor("input_209_cast_fp16")]; tensor hidden_states_91_axes_0 = const()[name = tensor("hidden_states_91_axes_0"), val = tensor([-1])]; tensor encoder_layers_15_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_15_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(457201024)))]; tensor encoder_layers_15_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_15_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(457203392)))]; tensor hidden_states_91_cast_fp16 = layer_norm(axes = hidden_states_91_axes_0, beta = encoder_layers_15_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_15_layer_norm1_weight_to_fp16, x = input_209_cast_fp16)[name = tensor("hidden_states_91_cast_fp16")]; tensor encoder_layers_15_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_15_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(457205760)))]; tensor encoder_layers_15_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_15_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(459860032)))]; tensor linear_90_cast_fp16 = linear(bias = encoder_layers_15_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_15_self_attn_q_proj_weight_to_fp16, x = hidden_states_91_cast_fp16)[name = tensor("linear_90_cast_fp16")]; tensor encoder_layers_15_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_15_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(459862400)))]; tensor encoder_layers_15_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_15_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(462516672)))]; tensor linear_91_cast_fp16 = linear(bias = encoder_layers_15_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_15_self_attn_k_proj_weight_to_fp16, x = hidden_states_91_cast_fp16)[name = tensor("linear_91_cast_fp16")]; tensor encoder_layers_15_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_15_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(462519040)))]; tensor encoder_layers_15_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_15_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(465173312)))]; tensor linear_92_cast_fp16 = linear(bias = encoder_layers_15_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_15_self_attn_v_proj_weight_to_fp16, x = hidden_states_91_cast_fp16)[name = tensor("linear_92_cast_fp16")]; tensor var_1026 = const()[name = tensor("op_1026"), val = tensor([1, 1024, 16, 72])]; tensor var_1027_cast_fp16 = reshape(shape = var_1026, x = linear_90_cast_fp16)[name = tensor("op_1027_cast_fp16")]; tensor var_1029 = const()[name = tensor("op_1029"), val = tensor([1, 1024, 16, 72])]; tensor var_1030_cast_fp16 = reshape(shape = var_1029, x = linear_91_cast_fp16)[name = tensor("op_1030_cast_fp16")]; tensor var_1032 = const()[name = tensor("op_1032"), val = tensor([1, 1024, 16, 72])]; tensor var_1033_cast_fp16 = reshape(shape = var_1032, x = linear_92_cast_fp16)[name = tensor("op_1033_cast_fp16")]; tensor value_states_63_perm_0 = const()[name = tensor("value_states_63_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_1036_transpose_x_0 = const()[name = tensor("op_1036_transpose_x_0"), val = tensor(false)]; tensor var_1036_transpose_y_0 = const()[name = tensor("op_1036_transpose_y_0"), val = tensor(false)]; tensor transpose_111_perm_0 = const()[name = tensor("transpose_111_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_112_perm_0 = const()[name = tensor("transpose_112_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_112 = transpose(perm = transpose_112_perm_0, x = var_1030_cast_fp16)[name = tensor("transpose_180")]; tensor transpose_111 = transpose(perm = transpose_111_perm_0, x = var_1027_cast_fp16)[name = tensor("transpose_181")]; tensor var_1036_cast_fp16 = matmul(transpose_x = var_1036_transpose_x_0, transpose_y = var_1036_transpose_y_0, x = transpose_111, y = transpose_112)[name = tensor("op_1036_cast_fp16")]; tensor var_1037_to_fp16 = const()[name = tensor("op_1037_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_61_cast_fp16 = mul(x = var_1036_cast_fp16, y = var_1037_to_fp16)[name = tensor("attn_weights_61_cast_fp16")]; tensor var_1039_cast_fp16 = softmax(axis = var_11, x = attn_weights_61_cast_fp16)[name = tensor("op_1039_cast_fp16")]; tensor attn_output_61_transpose_x_0 = const()[name = tensor("attn_output_61_transpose_x_0"), val = tensor(false)]; tensor attn_output_61_transpose_y_0 = const()[name = tensor("attn_output_61_transpose_y_0"), val = tensor(false)]; tensor value_states_63_cast_fp16 = transpose(perm = value_states_63_perm_0, x = var_1033_cast_fp16)[name = tensor("transpose_182")]; tensor attn_output_61_cast_fp16 = matmul(transpose_x = attn_output_61_transpose_x_0, transpose_y = attn_output_61_transpose_y_0, x = var_1039_cast_fp16, y = value_states_63_cast_fp16)[name = tensor("attn_output_61_cast_fp16")]; tensor var_1043_perm_0 = const()[name = tensor("op_1043_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_1045 = const()[name = tensor("op_1045"), val = tensor([1, 1024, 1152])]; tensor var_1043_cast_fp16 = transpose(perm = var_1043_perm_0, x = attn_output_61_cast_fp16)[name = tensor("transpose_179")]; tensor input_213_cast_fp16 = reshape(shape = var_1045, x = var_1043_cast_fp16)[name = tensor("input_213_cast_fp16")]; tensor encoder_layers_15_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_15_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(465175680)))]; tensor encoder_layers_15_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_15_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(467829952)))]; tensor linear_93_cast_fp16 = linear(bias = encoder_layers_15_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_15_self_attn_out_proj_weight_to_fp16, x = input_213_cast_fp16)[name = tensor("linear_93_cast_fp16")]; tensor input_215_cast_fp16 = add(x = input_209_cast_fp16, y = linear_93_cast_fp16)[name = tensor("input_215_cast_fp16")]; tensor input_217_axes_0 = const()[name = tensor("input_217_axes_0"), val = tensor([-1])]; tensor encoder_layers_15_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_15_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(467832320)))]; tensor encoder_layers_15_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_15_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(467834688)))]; tensor input_217_cast_fp16 = layer_norm(axes = input_217_axes_0, beta = encoder_layers_15_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_15_layer_norm2_weight_to_fp16, x = input_215_cast_fp16)[name = tensor("input_217_cast_fp16")]; tensor encoder_layers_15_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_15_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(467837056)))]; tensor encoder_layers_15_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_15_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(477753536)))]; tensor linear_94_cast_fp16 = linear(bias = encoder_layers_15_mlp_fc1_bias_to_fp16, weight = encoder_layers_15_mlp_fc1_weight_to_fp16, x = input_217_cast_fp16)[name = tensor("linear_94_cast_fp16")]; tensor input_221_mode_0 = const()[name = tensor("input_221_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_221_cast_fp16 = gelu(mode = input_221_mode_0, x = linear_94_cast_fp16)[name = tensor("input_221_cast_fp16")]; tensor encoder_layers_15_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_15_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(477762240)))]; tensor encoder_layers_15_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_15_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(487678720)))]; tensor linear_95_cast_fp16 = linear(bias = encoder_layers_15_mlp_fc2_bias_to_fp16, weight = encoder_layers_15_mlp_fc2_weight_to_fp16, x = input_221_cast_fp16)[name = tensor("linear_95_cast_fp16")]; tensor input_223_cast_fp16 = add(x = input_215_cast_fp16, y = linear_95_cast_fp16)[name = tensor("input_223_cast_fp16")]; tensor hidden_states_97_axes_0 = const()[name = tensor("hidden_states_97_axes_0"), val = tensor([-1])]; tensor encoder_layers_16_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_16_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(487681088)))]; tensor encoder_layers_16_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_16_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(487683456)))]; tensor hidden_states_97_cast_fp16 = layer_norm(axes = hidden_states_97_axes_0, beta = encoder_layers_16_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_16_layer_norm1_weight_to_fp16, x = input_223_cast_fp16)[name = tensor("hidden_states_97_cast_fp16")]; tensor encoder_layers_16_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_16_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(487685824)))]; tensor encoder_layers_16_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_16_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(490340096)))]; tensor linear_96_cast_fp16 = linear(bias = encoder_layers_16_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_16_self_attn_q_proj_weight_to_fp16, x = hidden_states_97_cast_fp16)[name = tensor("linear_96_cast_fp16")]; tensor encoder_layers_16_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_16_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(490342464)))]; tensor encoder_layers_16_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_16_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(492996736)))]; tensor linear_97_cast_fp16 = linear(bias = encoder_layers_16_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_16_self_attn_k_proj_weight_to_fp16, x = hidden_states_97_cast_fp16)[name = tensor("linear_97_cast_fp16")]; tensor encoder_layers_16_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_16_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(492999104)))]; tensor encoder_layers_16_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_16_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(495653376)))]; tensor linear_98_cast_fp16 = linear(bias = encoder_layers_16_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_16_self_attn_v_proj_weight_to_fp16, x = hidden_states_97_cast_fp16)[name = tensor("linear_98_cast_fp16")]; tensor var_1088 = const()[name = tensor("op_1088"), val = tensor([1, 1024, 16, 72])]; tensor var_1089_cast_fp16 = reshape(shape = var_1088, x = linear_96_cast_fp16)[name = tensor("op_1089_cast_fp16")]; tensor var_1091 = const()[name = tensor("op_1091"), val = tensor([1, 1024, 16, 72])]; tensor var_1092_cast_fp16 = reshape(shape = var_1091, x = linear_97_cast_fp16)[name = tensor("op_1092_cast_fp16")]; tensor var_1094 = const()[name = tensor("op_1094"), val = tensor([1, 1024, 16, 72])]; tensor var_1095_cast_fp16 = reshape(shape = var_1094, x = linear_98_cast_fp16)[name = tensor("op_1095_cast_fp16")]; tensor value_states_67_perm_0 = const()[name = tensor("value_states_67_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_1098_transpose_x_0 = const()[name = tensor("op_1098_transpose_x_0"), val = tensor(false)]; tensor var_1098_transpose_y_0 = const()[name = tensor("op_1098_transpose_y_0"), val = tensor(false)]; tensor transpose_113_perm_0 = const()[name = tensor("transpose_113_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_114_perm_0 = const()[name = tensor("transpose_114_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_114 = transpose(perm = transpose_114_perm_0, x = var_1092_cast_fp16)[name = tensor("transpose_176")]; tensor transpose_113 = transpose(perm = transpose_113_perm_0, x = var_1089_cast_fp16)[name = tensor("transpose_177")]; tensor var_1098_cast_fp16 = matmul(transpose_x = var_1098_transpose_x_0, transpose_y = var_1098_transpose_y_0, x = transpose_113, y = transpose_114)[name = tensor("op_1098_cast_fp16")]; tensor var_1099_to_fp16 = const()[name = tensor("op_1099_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_65_cast_fp16 = mul(x = var_1098_cast_fp16, y = var_1099_to_fp16)[name = tensor("attn_weights_65_cast_fp16")]; tensor var_1101_cast_fp16 = softmax(axis = var_11, x = attn_weights_65_cast_fp16)[name = tensor("op_1101_cast_fp16")]; tensor attn_output_65_transpose_x_0 = const()[name = tensor("attn_output_65_transpose_x_0"), val = tensor(false)]; tensor attn_output_65_transpose_y_0 = const()[name = tensor("attn_output_65_transpose_y_0"), val = tensor(false)]; tensor value_states_67_cast_fp16 = transpose(perm = value_states_67_perm_0, x = var_1095_cast_fp16)[name = tensor("transpose_178")]; tensor attn_output_65_cast_fp16 = matmul(transpose_x = attn_output_65_transpose_x_0, transpose_y = attn_output_65_transpose_y_0, x = var_1101_cast_fp16, y = value_states_67_cast_fp16)[name = tensor("attn_output_65_cast_fp16")]; tensor var_1105_perm_0 = const()[name = tensor("op_1105_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_1107 = const()[name = tensor("op_1107"), val = tensor([1, 1024, 1152])]; tensor var_1105_cast_fp16 = transpose(perm = var_1105_perm_0, x = attn_output_65_cast_fp16)[name = tensor("transpose_175")]; tensor input_227_cast_fp16 = reshape(shape = var_1107, x = var_1105_cast_fp16)[name = tensor("input_227_cast_fp16")]; tensor encoder_layers_16_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_16_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(495655744)))]; tensor encoder_layers_16_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_16_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(498310016)))]; tensor linear_99_cast_fp16 = linear(bias = encoder_layers_16_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_16_self_attn_out_proj_weight_to_fp16, x = input_227_cast_fp16)[name = tensor("linear_99_cast_fp16")]; tensor input_229_cast_fp16 = add(x = input_223_cast_fp16, y = linear_99_cast_fp16)[name = tensor("input_229_cast_fp16")]; tensor input_231_axes_0 = const()[name = tensor("input_231_axes_0"), val = tensor([-1])]; tensor encoder_layers_16_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_16_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(498312384)))]; tensor encoder_layers_16_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_16_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(498314752)))]; tensor input_231_cast_fp16 = layer_norm(axes = input_231_axes_0, beta = encoder_layers_16_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_16_layer_norm2_weight_to_fp16, x = input_229_cast_fp16)[name = tensor("input_231_cast_fp16")]; tensor encoder_layers_16_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_16_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(498317120)))]; tensor encoder_layers_16_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_16_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(508233600)))]; tensor linear_100_cast_fp16 = linear(bias = encoder_layers_16_mlp_fc1_bias_to_fp16, weight = encoder_layers_16_mlp_fc1_weight_to_fp16, x = input_231_cast_fp16)[name = tensor("linear_100_cast_fp16")]; tensor input_235_mode_0 = const()[name = tensor("input_235_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_235_cast_fp16 = gelu(mode = input_235_mode_0, x = linear_100_cast_fp16)[name = tensor("input_235_cast_fp16")]; tensor encoder_layers_16_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_16_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(508242304)))]; tensor encoder_layers_16_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_16_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(518158784)))]; tensor linear_101_cast_fp16 = linear(bias = encoder_layers_16_mlp_fc2_bias_to_fp16, weight = encoder_layers_16_mlp_fc2_weight_to_fp16, x = input_235_cast_fp16)[name = tensor("linear_101_cast_fp16")]; tensor input_237_cast_fp16 = add(x = input_229_cast_fp16, y = linear_101_cast_fp16)[name = tensor("input_237_cast_fp16")]; tensor hidden_states_103_axes_0 = const()[name = tensor("hidden_states_103_axes_0"), val = tensor([-1])]; tensor encoder_layers_17_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_17_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(518161152)))]; tensor encoder_layers_17_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_17_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(518163520)))]; tensor hidden_states_103_cast_fp16 = layer_norm(axes = hidden_states_103_axes_0, beta = encoder_layers_17_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_17_layer_norm1_weight_to_fp16, x = input_237_cast_fp16)[name = tensor("hidden_states_103_cast_fp16")]; tensor encoder_layers_17_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_17_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(518165888)))]; tensor encoder_layers_17_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_17_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(520820160)))]; tensor linear_102_cast_fp16 = linear(bias = encoder_layers_17_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_17_self_attn_q_proj_weight_to_fp16, x = hidden_states_103_cast_fp16)[name = tensor("linear_102_cast_fp16")]; tensor encoder_layers_17_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_17_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(520822528)))]; tensor encoder_layers_17_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_17_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(523476800)))]; tensor linear_103_cast_fp16 = linear(bias = encoder_layers_17_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_17_self_attn_k_proj_weight_to_fp16, x = hidden_states_103_cast_fp16)[name = tensor("linear_103_cast_fp16")]; tensor encoder_layers_17_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_17_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(523479168)))]; tensor encoder_layers_17_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_17_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(526133440)))]; tensor linear_104_cast_fp16 = linear(bias = encoder_layers_17_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_17_self_attn_v_proj_weight_to_fp16, x = hidden_states_103_cast_fp16)[name = tensor("linear_104_cast_fp16")]; tensor var_1150 = const()[name = tensor("op_1150"), val = tensor([1, 1024, 16, 72])]; tensor var_1151_cast_fp16 = reshape(shape = var_1150, x = linear_102_cast_fp16)[name = tensor("op_1151_cast_fp16")]; tensor var_1153 = const()[name = tensor("op_1153"), val = tensor([1, 1024, 16, 72])]; tensor var_1154_cast_fp16 = reshape(shape = var_1153, x = linear_103_cast_fp16)[name = tensor("op_1154_cast_fp16")]; tensor var_1156 = const()[name = tensor("op_1156"), val = tensor([1, 1024, 16, 72])]; tensor var_1157_cast_fp16 = reshape(shape = var_1156, x = linear_104_cast_fp16)[name = tensor("op_1157_cast_fp16")]; tensor value_states_71_perm_0 = const()[name = tensor("value_states_71_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_1160_transpose_x_0 = const()[name = tensor("op_1160_transpose_x_0"), val = tensor(false)]; tensor var_1160_transpose_y_0 = const()[name = tensor("op_1160_transpose_y_0"), val = tensor(false)]; tensor transpose_115_perm_0 = const()[name = tensor("transpose_115_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_116_perm_0 = const()[name = tensor("transpose_116_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_116 = transpose(perm = transpose_116_perm_0, x = var_1154_cast_fp16)[name = tensor("transpose_172")]; tensor transpose_115 = transpose(perm = transpose_115_perm_0, x = var_1151_cast_fp16)[name = tensor("transpose_173")]; tensor var_1160_cast_fp16 = matmul(transpose_x = var_1160_transpose_x_0, transpose_y = var_1160_transpose_y_0, x = transpose_115, y = transpose_116)[name = tensor("op_1160_cast_fp16")]; tensor var_1161_to_fp16 = const()[name = tensor("op_1161_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_69_cast_fp16 = mul(x = var_1160_cast_fp16, y = var_1161_to_fp16)[name = tensor("attn_weights_69_cast_fp16")]; tensor var_1163_cast_fp16 = softmax(axis = var_11, x = attn_weights_69_cast_fp16)[name = tensor("op_1163_cast_fp16")]; tensor attn_output_69_transpose_x_0 = const()[name = tensor("attn_output_69_transpose_x_0"), val = tensor(false)]; tensor attn_output_69_transpose_y_0 = const()[name = tensor("attn_output_69_transpose_y_0"), val = tensor(false)]; tensor value_states_71_cast_fp16 = transpose(perm = value_states_71_perm_0, x = var_1157_cast_fp16)[name = tensor("transpose_174")]; tensor attn_output_69_cast_fp16 = matmul(transpose_x = attn_output_69_transpose_x_0, transpose_y = attn_output_69_transpose_y_0, x = var_1163_cast_fp16, y = value_states_71_cast_fp16)[name = tensor("attn_output_69_cast_fp16")]; tensor var_1167_perm_0 = const()[name = tensor("op_1167_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_1169 = const()[name = tensor("op_1169"), val = tensor([1, 1024, 1152])]; tensor var_1167_cast_fp16 = transpose(perm = var_1167_perm_0, x = attn_output_69_cast_fp16)[name = tensor("transpose_171")]; tensor input_241_cast_fp16 = reshape(shape = var_1169, x = var_1167_cast_fp16)[name = tensor("input_241_cast_fp16")]; tensor encoder_layers_17_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_17_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(526135808)))]; tensor encoder_layers_17_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_17_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(528790080)))]; tensor linear_105_cast_fp16 = linear(bias = encoder_layers_17_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_17_self_attn_out_proj_weight_to_fp16, x = input_241_cast_fp16)[name = tensor("linear_105_cast_fp16")]; tensor input_243_cast_fp16 = add(x = input_237_cast_fp16, y = linear_105_cast_fp16)[name = tensor("input_243_cast_fp16")]; tensor input_245_axes_0 = const()[name = tensor("input_245_axes_0"), val = tensor([-1])]; tensor encoder_layers_17_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_17_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(528792448)))]; tensor encoder_layers_17_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_17_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(528794816)))]; tensor input_245_cast_fp16 = layer_norm(axes = input_245_axes_0, beta = encoder_layers_17_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_17_layer_norm2_weight_to_fp16, x = input_243_cast_fp16)[name = tensor("input_245_cast_fp16")]; tensor encoder_layers_17_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_17_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(528797184)))]; tensor encoder_layers_17_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_17_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(538713664)))]; tensor linear_106_cast_fp16 = linear(bias = encoder_layers_17_mlp_fc1_bias_to_fp16, weight = encoder_layers_17_mlp_fc1_weight_to_fp16, x = input_245_cast_fp16)[name = tensor("linear_106_cast_fp16")]; tensor input_249_mode_0 = const()[name = tensor("input_249_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_249_cast_fp16 = gelu(mode = input_249_mode_0, x = linear_106_cast_fp16)[name = tensor("input_249_cast_fp16")]; tensor encoder_layers_17_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_17_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(538722368)))]; tensor encoder_layers_17_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_17_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(548638848)))]; tensor linear_107_cast_fp16 = linear(bias = encoder_layers_17_mlp_fc2_bias_to_fp16, weight = encoder_layers_17_mlp_fc2_weight_to_fp16, x = input_249_cast_fp16)[name = tensor("linear_107_cast_fp16")]; tensor input_251_cast_fp16 = add(x = input_243_cast_fp16, y = linear_107_cast_fp16)[name = tensor("input_251_cast_fp16")]; tensor hidden_states_109_axes_0 = const()[name = tensor("hidden_states_109_axes_0"), val = tensor([-1])]; tensor encoder_layers_18_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_18_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(548641216)))]; tensor encoder_layers_18_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_18_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(548643584)))]; tensor hidden_states_109_cast_fp16 = layer_norm(axes = hidden_states_109_axes_0, beta = encoder_layers_18_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_18_layer_norm1_weight_to_fp16, x = input_251_cast_fp16)[name = tensor("hidden_states_109_cast_fp16")]; tensor encoder_layers_18_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_18_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(548645952)))]; tensor encoder_layers_18_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_18_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(551300224)))]; tensor linear_108_cast_fp16 = linear(bias = encoder_layers_18_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_18_self_attn_q_proj_weight_to_fp16, x = hidden_states_109_cast_fp16)[name = tensor("linear_108_cast_fp16")]; tensor encoder_layers_18_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_18_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(551302592)))]; tensor encoder_layers_18_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_18_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(553956864)))]; tensor linear_109_cast_fp16 = linear(bias = encoder_layers_18_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_18_self_attn_k_proj_weight_to_fp16, x = hidden_states_109_cast_fp16)[name = tensor("linear_109_cast_fp16")]; tensor encoder_layers_18_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_18_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(553959232)))]; tensor encoder_layers_18_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_18_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(556613504)))]; tensor linear_110_cast_fp16 = linear(bias = encoder_layers_18_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_18_self_attn_v_proj_weight_to_fp16, x = hidden_states_109_cast_fp16)[name = tensor("linear_110_cast_fp16")]; tensor var_1212 = const()[name = tensor("op_1212"), val = tensor([1, 1024, 16, 72])]; tensor var_1213_cast_fp16 = reshape(shape = var_1212, x = linear_108_cast_fp16)[name = tensor("op_1213_cast_fp16")]; tensor var_1215 = const()[name = tensor("op_1215"), val = tensor([1, 1024, 16, 72])]; tensor var_1216_cast_fp16 = reshape(shape = var_1215, x = linear_109_cast_fp16)[name = tensor("op_1216_cast_fp16")]; tensor var_1218 = const()[name = tensor("op_1218"), val = tensor([1, 1024, 16, 72])]; tensor var_1219_cast_fp16 = reshape(shape = var_1218, x = linear_110_cast_fp16)[name = tensor("op_1219_cast_fp16")]; tensor value_states_75_perm_0 = const()[name = tensor("value_states_75_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_1222_transpose_x_0 = const()[name = tensor("op_1222_transpose_x_0"), val = tensor(false)]; tensor var_1222_transpose_y_0 = const()[name = tensor("op_1222_transpose_y_0"), val = tensor(false)]; tensor transpose_117_perm_0 = const()[name = tensor("transpose_117_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_118_perm_0 = const()[name = tensor("transpose_118_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_118 = transpose(perm = transpose_118_perm_0, x = var_1216_cast_fp16)[name = tensor("transpose_168")]; tensor transpose_117 = transpose(perm = transpose_117_perm_0, x = var_1213_cast_fp16)[name = tensor("transpose_169")]; tensor var_1222_cast_fp16 = matmul(transpose_x = var_1222_transpose_x_0, transpose_y = var_1222_transpose_y_0, x = transpose_117, y = transpose_118)[name = tensor("op_1222_cast_fp16")]; tensor var_1223_to_fp16 = const()[name = tensor("op_1223_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_73_cast_fp16 = mul(x = var_1222_cast_fp16, y = var_1223_to_fp16)[name = tensor("attn_weights_73_cast_fp16")]; tensor var_1225_cast_fp16 = softmax(axis = var_11, x = attn_weights_73_cast_fp16)[name = tensor("op_1225_cast_fp16")]; tensor attn_output_73_transpose_x_0 = const()[name = tensor("attn_output_73_transpose_x_0"), val = tensor(false)]; tensor attn_output_73_transpose_y_0 = const()[name = tensor("attn_output_73_transpose_y_0"), val = tensor(false)]; tensor value_states_75_cast_fp16 = transpose(perm = value_states_75_perm_0, x = var_1219_cast_fp16)[name = tensor("transpose_170")]; tensor attn_output_73_cast_fp16 = matmul(transpose_x = attn_output_73_transpose_x_0, transpose_y = attn_output_73_transpose_y_0, x = var_1225_cast_fp16, y = value_states_75_cast_fp16)[name = tensor("attn_output_73_cast_fp16")]; tensor var_1229_perm_0 = const()[name = tensor("op_1229_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_1231 = const()[name = tensor("op_1231"), val = tensor([1, 1024, 1152])]; tensor var_1229_cast_fp16 = transpose(perm = var_1229_perm_0, x = attn_output_73_cast_fp16)[name = tensor("transpose_167")]; tensor input_255_cast_fp16 = reshape(shape = var_1231, x = var_1229_cast_fp16)[name = tensor("input_255_cast_fp16")]; tensor encoder_layers_18_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_18_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(556615872)))]; tensor encoder_layers_18_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_18_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(559270144)))]; tensor linear_111_cast_fp16 = linear(bias = encoder_layers_18_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_18_self_attn_out_proj_weight_to_fp16, x = input_255_cast_fp16)[name = tensor("linear_111_cast_fp16")]; tensor input_257_cast_fp16 = add(x = input_251_cast_fp16, y = linear_111_cast_fp16)[name = tensor("input_257_cast_fp16")]; tensor input_259_axes_0 = const()[name = tensor("input_259_axes_0"), val = tensor([-1])]; tensor encoder_layers_18_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_18_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(559272512)))]; tensor encoder_layers_18_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_18_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(559274880)))]; tensor input_259_cast_fp16 = layer_norm(axes = input_259_axes_0, beta = encoder_layers_18_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_18_layer_norm2_weight_to_fp16, x = input_257_cast_fp16)[name = tensor("input_259_cast_fp16")]; tensor encoder_layers_18_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_18_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(559277248)))]; tensor encoder_layers_18_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_18_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(569193728)))]; tensor linear_112_cast_fp16 = linear(bias = encoder_layers_18_mlp_fc1_bias_to_fp16, weight = encoder_layers_18_mlp_fc1_weight_to_fp16, x = input_259_cast_fp16)[name = tensor("linear_112_cast_fp16")]; tensor input_263_mode_0 = const()[name = tensor("input_263_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_263_cast_fp16 = gelu(mode = input_263_mode_0, x = linear_112_cast_fp16)[name = tensor("input_263_cast_fp16")]; tensor encoder_layers_18_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_18_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(569202432)))]; tensor encoder_layers_18_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_18_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(579118912)))]; tensor linear_113_cast_fp16 = linear(bias = encoder_layers_18_mlp_fc2_bias_to_fp16, weight = encoder_layers_18_mlp_fc2_weight_to_fp16, x = input_263_cast_fp16)[name = tensor("linear_113_cast_fp16")]; tensor input_265_cast_fp16 = add(x = input_257_cast_fp16, y = linear_113_cast_fp16)[name = tensor("input_265_cast_fp16")]; tensor hidden_states_115_axes_0 = const()[name = tensor("hidden_states_115_axes_0"), val = tensor([-1])]; tensor encoder_layers_19_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_19_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(579121280)))]; tensor encoder_layers_19_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_19_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(579123648)))]; tensor hidden_states_115_cast_fp16 = layer_norm(axes = hidden_states_115_axes_0, beta = encoder_layers_19_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_19_layer_norm1_weight_to_fp16, x = input_265_cast_fp16)[name = tensor("hidden_states_115_cast_fp16")]; tensor encoder_layers_19_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_19_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(579126016)))]; tensor encoder_layers_19_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_19_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(581780288)))]; tensor linear_114_cast_fp16 = linear(bias = encoder_layers_19_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_19_self_attn_q_proj_weight_to_fp16, x = hidden_states_115_cast_fp16)[name = tensor("linear_114_cast_fp16")]; tensor encoder_layers_19_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_19_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(581782656)))]; tensor encoder_layers_19_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_19_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(584436928)))]; tensor linear_115_cast_fp16 = linear(bias = encoder_layers_19_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_19_self_attn_k_proj_weight_to_fp16, x = hidden_states_115_cast_fp16)[name = tensor("linear_115_cast_fp16")]; tensor encoder_layers_19_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_19_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(584439296)))]; tensor encoder_layers_19_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_19_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(587093568)))]; tensor linear_116_cast_fp16 = linear(bias = encoder_layers_19_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_19_self_attn_v_proj_weight_to_fp16, x = hidden_states_115_cast_fp16)[name = tensor("linear_116_cast_fp16")]; tensor var_1274 = const()[name = tensor("op_1274"), val = tensor([1, 1024, 16, 72])]; tensor var_1275_cast_fp16 = reshape(shape = var_1274, x = linear_114_cast_fp16)[name = tensor("op_1275_cast_fp16")]; tensor var_1277 = const()[name = tensor("op_1277"), val = tensor([1, 1024, 16, 72])]; tensor var_1278_cast_fp16 = reshape(shape = var_1277, x = linear_115_cast_fp16)[name = tensor("op_1278_cast_fp16")]; tensor var_1280 = const()[name = tensor("op_1280"), val = tensor([1, 1024, 16, 72])]; tensor var_1281_cast_fp16 = reshape(shape = var_1280, x = linear_116_cast_fp16)[name = tensor("op_1281_cast_fp16")]; tensor value_states_79_perm_0 = const()[name = tensor("value_states_79_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_1284_transpose_x_0 = const()[name = tensor("op_1284_transpose_x_0"), val = tensor(false)]; tensor var_1284_transpose_y_0 = const()[name = tensor("op_1284_transpose_y_0"), val = tensor(false)]; tensor transpose_119_perm_0 = const()[name = tensor("transpose_119_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_120_perm_0 = const()[name = tensor("transpose_120_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_120 = transpose(perm = transpose_120_perm_0, x = var_1278_cast_fp16)[name = tensor("transpose_164")]; tensor transpose_119 = transpose(perm = transpose_119_perm_0, x = var_1275_cast_fp16)[name = tensor("transpose_165")]; tensor var_1284_cast_fp16 = matmul(transpose_x = var_1284_transpose_x_0, transpose_y = var_1284_transpose_y_0, x = transpose_119, y = transpose_120)[name = tensor("op_1284_cast_fp16")]; tensor var_1285_to_fp16 = const()[name = tensor("op_1285_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_77_cast_fp16 = mul(x = var_1284_cast_fp16, y = var_1285_to_fp16)[name = tensor("attn_weights_77_cast_fp16")]; tensor var_1287_cast_fp16 = softmax(axis = var_11, x = attn_weights_77_cast_fp16)[name = tensor("op_1287_cast_fp16")]; tensor attn_output_77_transpose_x_0 = const()[name = tensor("attn_output_77_transpose_x_0"), val = tensor(false)]; tensor attn_output_77_transpose_y_0 = const()[name = tensor("attn_output_77_transpose_y_0"), val = tensor(false)]; tensor value_states_79_cast_fp16 = transpose(perm = value_states_79_perm_0, x = var_1281_cast_fp16)[name = tensor("transpose_166")]; tensor attn_output_77_cast_fp16 = matmul(transpose_x = attn_output_77_transpose_x_0, transpose_y = attn_output_77_transpose_y_0, x = var_1287_cast_fp16, y = value_states_79_cast_fp16)[name = tensor("attn_output_77_cast_fp16")]; tensor var_1291_perm_0 = const()[name = tensor("op_1291_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_1293 = const()[name = tensor("op_1293"), val = tensor([1, 1024, 1152])]; tensor var_1291_cast_fp16 = transpose(perm = var_1291_perm_0, x = attn_output_77_cast_fp16)[name = tensor("transpose_163")]; tensor input_269_cast_fp16 = reshape(shape = var_1293, x = var_1291_cast_fp16)[name = tensor("input_269_cast_fp16")]; tensor encoder_layers_19_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_19_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(587095936)))]; tensor encoder_layers_19_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_19_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(589750208)))]; tensor linear_117_cast_fp16 = linear(bias = encoder_layers_19_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_19_self_attn_out_proj_weight_to_fp16, x = input_269_cast_fp16)[name = tensor("linear_117_cast_fp16")]; tensor input_271_cast_fp16 = add(x = input_265_cast_fp16, y = linear_117_cast_fp16)[name = tensor("input_271_cast_fp16")]; tensor input_273_axes_0 = const()[name = tensor("input_273_axes_0"), val = tensor([-1])]; tensor encoder_layers_19_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_19_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(589752576)))]; tensor encoder_layers_19_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_19_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(589754944)))]; tensor input_273_cast_fp16 = layer_norm(axes = input_273_axes_0, beta = encoder_layers_19_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_19_layer_norm2_weight_to_fp16, x = input_271_cast_fp16)[name = tensor("input_273_cast_fp16")]; tensor encoder_layers_19_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_19_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(589757312)))]; tensor encoder_layers_19_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_19_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(599673792)))]; tensor linear_118_cast_fp16 = linear(bias = encoder_layers_19_mlp_fc1_bias_to_fp16, weight = encoder_layers_19_mlp_fc1_weight_to_fp16, x = input_273_cast_fp16)[name = tensor("linear_118_cast_fp16")]; tensor input_277_mode_0 = const()[name = tensor("input_277_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_277_cast_fp16 = gelu(mode = input_277_mode_0, x = linear_118_cast_fp16)[name = tensor("input_277_cast_fp16")]; tensor encoder_layers_19_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_19_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(599682496)))]; tensor encoder_layers_19_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_19_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(609598976)))]; tensor linear_119_cast_fp16 = linear(bias = encoder_layers_19_mlp_fc2_bias_to_fp16, weight = encoder_layers_19_mlp_fc2_weight_to_fp16, x = input_277_cast_fp16)[name = tensor("linear_119_cast_fp16")]; tensor input_279_cast_fp16 = add(x = input_271_cast_fp16, y = linear_119_cast_fp16)[name = tensor("input_279_cast_fp16")]; tensor hidden_states_121_axes_0 = const()[name = tensor("hidden_states_121_axes_0"), val = tensor([-1])]; tensor encoder_layers_20_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_20_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(609601344)))]; tensor encoder_layers_20_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_20_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(609603712)))]; tensor hidden_states_121_cast_fp16 = layer_norm(axes = hidden_states_121_axes_0, beta = encoder_layers_20_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_20_layer_norm1_weight_to_fp16, x = input_279_cast_fp16)[name = tensor("hidden_states_121_cast_fp16")]; tensor encoder_layers_20_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_20_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(609606080)))]; tensor encoder_layers_20_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_20_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(612260352)))]; tensor linear_120_cast_fp16 = linear(bias = encoder_layers_20_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_20_self_attn_q_proj_weight_to_fp16, x = hidden_states_121_cast_fp16)[name = tensor("linear_120_cast_fp16")]; tensor encoder_layers_20_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_20_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(612262720)))]; tensor encoder_layers_20_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_20_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(614916992)))]; tensor linear_121_cast_fp16 = linear(bias = encoder_layers_20_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_20_self_attn_k_proj_weight_to_fp16, x = hidden_states_121_cast_fp16)[name = tensor("linear_121_cast_fp16")]; tensor encoder_layers_20_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_20_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(614919360)))]; tensor encoder_layers_20_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_20_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(617573632)))]; tensor linear_122_cast_fp16 = linear(bias = encoder_layers_20_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_20_self_attn_v_proj_weight_to_fp16, x = hidden_states_121_cast_fp16)[name = tensor("linear_122_cast_fp16")]; tensor var_1336 = const()[name = tensor("op_1336"), val = tensor([1, 1024, 16, 72])]; tensor var_1337_cast_fp16 = reshape(shape = var_1336, x = linear_120_cast_fp16)[name = tensor("op_1337_cast_fp16")]; tensor var_1339 = const()[name = tensor("op_1339"), val = tensor([1, 1024, 16, 72])]; tensor var_1340_cast_fp16 = reshape(shape = var_1339, x = linear_121_cast_fp16)[name = tensor("op_1340_cast_fp16")]; tensor var_1342 = const()[name = tensor("op_1342"), val = tensor([1, 1024, 16, 72])]; tensor var_1343_cast_fp16 = reshape(shape = var_1342, x = linear_122_cast_fp16)[name = tensor("op_1343_cast_fp16")]; tensor value_states_83_perm_0 = const()[name = tensor("value_states_83_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_1346_transpose_x_0 = const()[name = tensor("op_1346_transpose_x_0"), val = tensor(false)]; tensor var_1346_transpose_y_0 = const()[name = tensor("op_1346_transpose_y_0"), val = tensor(false)]; tensor transpose_121_perm_0 = const()[name = tensor("transpose_121_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_122_perm_0 = const()[name = tensor("transpose_122_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_122 = transpose(perm = transpose_122_perm_0, x = var_1340_cast_fp16)[name = tensor("transpose_160")]; tensor transpose_121 = transpose(perm = transpose_121_perm_0, x = var_1337_cast_fp16)[name = tensor("transpose_161")]; tensor var_1346_cast_fp16 = matmul(transpose_x = var_1346_transpose_x_0, transpose_y = var_1346_transpose_y_0, x = transpose_121, y = transpose_122)[name = tensor("op_1346_cast_fp16")]; tensor var_1347_to_fp16 = const()[name = tensor("op_1347_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_81_cast_fp16 = mul(x = var_1346_cast_fp16, y = var_1347_to_fp16)[name = tensor("attn_weights_81_cast_fp16")]; tensor var_1349_cast_fp16 = softmax(axis = var_11, x = attn_weights_81_cast_fp16)[name = tensor("op_1349_cast_fp16")]; tensor attn_output_81_transpose_x_0 = const()[name = tensor("attn_output_81_transpose_x_0"), val = tensor(false)]; tensor attn_output_81_transpose_y_0 = const()[name = tensor("attn_output_81_transpose_y_0"), val = tensor(false)]; tensor value_states_83_cast_fp16 = transpose(perm = value_states_83_perm_0, x = var_1343_cast_fp16)[name = tensor("transpose_162")]; tensor attn_output_81_cast_fp16 = matmul(transpose_x = attn_output_81_transpose_x_0, transpose_y = attn_output_81_transpose_y_0, x = var_1349_cast_fp16, y = value_states_83_cast_fp16)[name = tensor("attn_output_81_cast_fp16")]; tensor var_1353_perm_0 = const()[name = tensor("op_1353_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_1355 = const()[name = tensor("op_1355"), val = tensor([1, 1024, 1152])]; tensor var_1353_cast_fp16 = transpose(perm = var_1353_perm_0, x = attn_output_81_cast_fp16)[name = tensor("transpose_159")]; tensor input_283_cast_fp16 = reshape(shape = var_1355, x = var_1353_cast_fp16)[name = tensor("input_283_cast_fp16")]; tensor encoder_layers_20_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_20_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(617576000)))]; tensor encoder_layers_20_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_20_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(620230272)))]; tensor linear_123_cast_fp16 = linear(bias = encoder_layers_20_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_20_self_attn_out_proj_weight_to_fp16, x = input_283_cast_fp16)[name = tensor("linear_123_cast_fp16")]; tensor input_285_cast_fp16 = add(x = input_279_cast_fp16, y = linear_123_cast_fp16)[name = tensor("input_285_cast_fp16")]; tensor input_287_axes_0 = const()[name = tensor("input_287_axes_0"), val = tensor([-1])]; tensor encoder_layers_20_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_20_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(620232640)))]; tensor encoder_layers_20_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_20_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(620235008)))]; tensor input_287_cast_fp16 = layer_norm(axes = input_287_axes_0, beta = encoder_layers_20_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_20_layer_norm2_weight_to_fp16, x = input_285_cast_fp16)[name = tensor("input_287_cast_fp16")]; tensor encoder_layers_20_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_20_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(620237376)))]; tensor encoder_layers_20_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_20_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(630153856)))]; tensor linear_124_cast_fp16 = linear(bias = encoder_layers_20_mlp_fc1_bias_to_fp16, weight = encoder_layers_20_mlp_fc1_weight_to_fp16, x = input_287_cast_fp16)[name = tensor("linear_124_cast_fp16")]; tensor input_291_mode_0 = const()[name = tensor("input_291_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_291_cast_fp16 = gelu(mode = input_291_mode_0, x = linear_124_cast_fp16)[name = tensor("input_291_cast_fp16")]; tensor encoder_layers_20_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_20_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(630162560)))]; tensor encoder_layers_20_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_20_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(640079040)))]; tensor linear_125_cast_fp16 = linear(bias = encoder_layers_20_mlp_fc2_bias_to_fp16, weight = encoder_layers_20_mlp_fc2_weight_to_fp16, x = input_291_cast_fp16)[name = tensor("linear_125_cast_fp16")]; tensor input_293_cast_fp16 = add(x = input_285_cast_fp16, y = linear_125_cast_fp16)[name = tensor("input_293_cast_fp16")]; tensor hidden_states_127_axes_0 = const()[name = tensor("hidden_states_127_axes_0"), val = tensor([-1])]; tensor encoder_layers_21_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_21_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(640081408)))]; tensor encoder_layers_21_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_21_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(640083776)))]; tensor hidden_states_127_cast_fp16 = layer_norm(axes = hidden_states_127_axes_0, beta = encoder_layers_21_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_21_layer_norm1_weight_to_fp16, x = input_293_cast_fp16)[name = tensor("hidden_states_127_cast_fp16")]; tensor encoder_layers_21_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_21_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(640086144)))]; tensor encoder_layers_21_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_21_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(642740416)))]; tensor linear_126_cast_fp16 = linear(bias = encoder_layers_21_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_21_self_attn_q_proj_weight_to_fp16, x = hidden_states_127_cast_fp16)[name = tensor("linear_126_cast_fp16")]; tensor encoder_layers_21_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_21_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(642742784)))]; tensor encoder_layers_21_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_21_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(645397056)))]; tensor linear_127_cast_fp16 = linear(bias = encoder_layers_21_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_21_self_attn_k_proj_weight_to_fp16, x = hidden_states_127_cast_fp16)[name = tensor("linear_127_cast_fp16")]; tensor encoder_layers_21_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_21_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(645399424)))]; tensor encoder_layers_21_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_21_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(648053696)))]; tensor linear_128_cast_fp16 = linear(bias = encoder_layers_21_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_21_self_attn_v_proj_weight_to_fp16, x = hidden_states_127_cast_fp16)[name = tensor("linear_128_cast_fp16")]; tensor var_1398 = const()[name = tensor("op_1398"), val = tensor([1, 1024, 16, 72])]; tensor var_1399_cast_fp16 = reshape(shape = var_1398, x = linear_126_cast_fp16)[name = tensor("op_1399_cast_fp16")]; tensor var_1401 = const()[name = tensor("op_1401"), val = tensor([1, 1024, 16, 72])]; tensor var_1402_cast_fp16 = reshape(shape = var_1401, x = linear_127_cast_fp16)[name = tensor("op_1402_cast_fp16")]; tensor var_1404 = const()[name = tensor("op_1404"), val = tensor([1, 1024, 16, 72])]; tensor var_1405_cast_fp16 = reshape(shape = var_1404, x = linear_128_cast_fp16)[name = tensor("op_1405_cast_fp16")]; tensor value_states_87_perm_0 = const()[name = tensor("value_states_87_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_1408_transpose_x_0 = const()[name = tensor("op_1408_transpose_x_0"), val = tensor(false)]; tensor var_1408_transpose_y_0 = const()[name = tensor("op_1408_transpose_y_0"), val = tensor(false)]; tensor transpose_123_perm_0 = const()[name = tensor("transpose_123_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_124_perm_0 = const()[name = tensor("transpose_124_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_124 = transpose(perm = transpose_124_perm_0, x = var_1402_cast_fp16)[name = tensor("transpose_156")]; tensor transpose_123 = transpose(perm = transpose_123_perm_0, x = var_1399_cast_fp16)[name = tensor("transpose_157")]; tensor var_1408_cast_fp16 = matmul(transpose_x = var_1408_transpose_x_0, transpose_y = var_1408_transpose_y_0, x = transpose_123, y = transpose_124)[name = tensor("op_1408_cast_fp16")]; tensor var_1409_to_fp16 = const()[name = tensor("op_1409_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_85_cast_fp16 = mul(x = var_1408_cast_fp16, y = var_1409_to_fp16)[name = tensor("attn_weights_85_cast_fp16")]; tensor var_1411_cast_fp16 = softmax(axis = var_11, x = attn_weights_85_cast_fp16)[name = tensor("op_1411_cast_fp16")]; tensor attn_output_85_transpose_x_0 = const()[name = tensor("attn_output_85_transpose_x_0"), val = tensor(false)]; tensor attn_output_85_transpose_y_0 = const()[name = tensor("attn_output_85_transpose_y_0"), val = tensor(false)]; tensor value_states_87_cast_fp16 = transpose(perm = value_states_87_perm_0, x = var_1405_cast_fp16)[name = tensor("transpose_158")]; tensor attn_output_85_cast_fp16 = matmul(transpose_x = attn_output_85_transpose_x_0, transpose_y = attn_output_85_transpose_y_0, x = var_1411_cast_fp16, y = value_states_87_cast_fp16)[name = tensor("attn_output_85_cast_fp16")]; tensor var_1415_perm_0 = const()[name = tensor("op_1415_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_1417 = const()[name = tensor("op_1417"), val = tensor([1, 1024, 1152])]; tensor var_1415_cast_fp16 = transpose(perm = var_1415_perm_0, x = attn_output_85_cast_fp16)[name = tensor("transpose_155")]; tensor input_297_cast_fp16 = reshape(shape = var_1417, x = var_1415_cast_fp16)[name = tensor("input_297_cast_fp16")]; tensor encoder_layers_21_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_21_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(648056064)))]; tensor encoder_layers_21_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_21_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(650710336)))]; tensor linear_129_cast_fp16 = linear(bias = encoder_layers_21_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_21_self_attn_out_proj_weight_to_fp16, x = input_297_cast_fp16)[name = tensor("linear_129_cast_fp16")]; tensor input_299_cast_fp16 = add(x = input_293_cast_fp16, y = linear_129_cast_fp16)[name = tensor("input_299_cast_fp16")]; tensor input_301_axes_0 = const()[name = tensor("input_301_axes_0"), val = tensor([-1])]; tensor encoder_layers_21_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_21_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(650712704)))]; tensor encoder_layers_21_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_21_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(650715072)))]; tensor input_301_cast_fp16 = layer_norm(axes = input_301_axes_0, beta = encoder_layers_21_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_21_layer_norm2_weight_to_fp16, x = input_299_cast_fp16)[name = tensor("input_301_cast_fp16")]; tensor encoder_layers_21_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_21_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(650717440)))]; tensor encoder_layers_21_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_21_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(660633920)))]; tensor linear_130_cast_fp16 = linear(bias = encoder_layers_21_mlp_fc1_bias_to_fp16, weight = encoder_layers_21_mlp_fc1_weight_to_fp16, x = input_301_cast_fp16)[name = tensor("linear_130_cast_fp16")]; tensor input_305_mode_0 = const()[name = tensor("input_305_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_305_cast_fp16 = gelu(mode = input_305_mode_0, x = linear_130_cast_fp16)[name = tensor("input_305_cast_fp16")]; tensor encoder_layers_21_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_21_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(660642624)))]; tensor encoder_layers_21_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_21_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(670559104)))]; tensor linear_131_cast_fp16 = linear(bias = encoder_layers_21_mlp_fc2_bias_to_fp16, weight = encoder_layers_21_mlp_fc2_weight_to_fp16, x = input_305_cast_fp16)[name = tensor("linear_131_cast_fp16")]; tensor input_307_cast_fp16 = add(x = input_299_cast_fp16, y = linear_131_cast_fp16)[name = tensor("input_307_cast_fp16")]; tensor hidden_states_133_axes_0 = const()[name = tensor("hidden_states_133_axes_0"), val = tensor([-1])]; tensor encoder_layers_22_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_22_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(670561472)))]; tensor encoder_layers_22_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_22_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(670563840)))]; tensor hidden_states_133_cast_fp16 = layer_norm(axes = hidden_states_133_axes_0, beta = encoder_layers_22_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_22_layer_norm1_weight_to_fp16, x = input_307_cast_fp16)[name = tensor("hidden_states_133_cast_fp16")]; tensor encoder_layers_22_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_22_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(670566208)))]; tensor encoder_layers_22_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_22_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(673220480)))]; tensor linear_132_cast_fp16 = linear(bias = encoder_layers_22_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_22_self_attn_q_proj_weight_to_fp16, x = hidden_states_133_cast_fp16)[name = tensor("linear_132_cast_fp16")]; tensor encoder_layers_22_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_22_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(673222848)))]; tensor encoder_layers_22_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_22_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(675877120)))]; tensor linear_133_cast_fp16 = linear(bias = encoder_layers_22_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_22_self_attn_k_proj_weight_to_fp16, x = hidden_states_133_cast_fp16)[name = tensor("linear_133_cast_fp16")]; tensor encoder_layers_22_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_22_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(675879488)))]; tensor encoder_layers_22_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_22_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(678533760)))]; tensor linear_134_cast_fp16 = linear(bias = encoder_layers_22_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_22_self_attn_v_proj_weight_to_fp16, x = hidden_states_133_cast_fp16)[name = tensor("linear_134_cast_fp16")]; tensor var_1460 = const()[name = tensor("op_1460"), val = tensor([1, 1024, 16, 72])]; tensor var_1461_cast_fp16 = reshape(shape = var_1460, x = linear_132_cast_fp16)[name = tensor("op_1461_cast_fp16")]; tensor var_1463 = const()[name = tensor("op_1463"), val = tensor([1, 1024, 16, 72])]; tensor var_1464_cast_fp16 = reshape(shape = var_1463, x = linear_133_cast_fp16)[name = tensor("op_1464_cast_fp16")]; tensor var_1466 = const()[name = tensor("op_1466"), val = tensor([1, 1024, 16, 72])]; tensor var_1467_cast_fp16 = reshape(shape = var_1466, x = linear_134_cast_fp16)[name = tensor("op_1467_cast_fp16")]; tensor value_states_91_perm_0 = const()[name = tensor("value_states_91_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_1470_transpose_x_0 = const()[name = tensor("op_1470_transpose_x_0"), val = tensor(false)]; tensor var_1470_transpose_y_0 = const()[name = tensor("op_1470_transpose_y_0"), val = tensor(false)]; tensor transpose_125_perm_0 = const()[name = tensor("transpose_125_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_126_perm_0 = const()[name = tensor("transpose_126_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_126 = transpose(perm = transpose_126_perm_0, x = var_1464_cast_fp16)[name = tensor("transpose_152")]; tensor transpose_125 = transpose(perm = transpose_125_perm_0, x = var_1461_cast_fp16)[name = tensor("transpose_153")]; tensor var_1470_cast_fp16 = matmul(transpose_x = var_1470_transpose_x_0, transpose_y = var_1470_transpose_y_0, x = transpose_125, y = transpose_126)[name = tensor("op_1470_cast_fp16")]; tensor var_1471_to_fp16 = const()[name = tensor("op_1471_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_89_cast_fp16 = mul(x = var_1470_cast_fp16, y = var_1471_to_fp16)[name = tensor("attn_weights_89_cast_fp16")]; tensor var_1473_cast_fp16 = softmax(axis = var_11, x = attn_weights_89_cast_fp16)[name = tensor("op_1473_cast_fp16")]; tensor attn_output_89_transpose_x_0 = const()[name = tensor("attn_output_89_transpose_x_0"), val = tensor(false)]; tensor attn_output_89_transpose_y_0 = const()[name = tensor("attn_output_89_transpose_y_0"), val = tensor(false)]; tensor value_states_91_cast_fp16 = transpose(perm = value_states_91_perm_0, x = var_1467_cast_fp16)[name = tensor("transpose_154")]; tensor attn_output_89_cast_fp16 = matmul(transpose_x = attn_output_89_transpose_x_0, transpose_y = attn_output_89_transpose_y_0, x = var_1473_cast_fp16, y = value_states_91_cast_fp16)[name = tensor("attn_output_89_cast_fp16")]; tensor var_1477_perm_0 = const()[name = tensor("op_1477_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_1479 = const()[name = tensor("op_1479"), val = tensor([1, 1024, 1152])]; tensor var_1477_cast_fp16 = transpose(perm = var_1477_perm_0, x = attn_output_89_cast_fp16)[name = tensor("transpose_151")]; tensor input_311_cast_fp16 = reshape(shape = var_1479, x = var_1477_cast_fp16)[name = tensor("input_311_cast_fp16")]; tensor encoder_layers_22_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_22_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(678536128)))]; tensor encoder_layers_22_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_22_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(681190400)))]; tensor linear_135_cast_fp16 = linear(bias = encoder_layers_22_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_22_self_attn_out_proj_weight_to_fp16, x = input_311_cast_fp16)[name = tensor("linear_135_cast_fp16")]; tensor input_313_cast_fp16 = add(x = input_307_cast_fp16, y = linear_135_cast_fp16)[name = tensor("input_313_cast_fp16")]; tensor input_315_axes_0 = const()[name = tensor("input_315_axes_0"), val = tensor([-1])]; tensor encoder_layers_22_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_22_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(681192768)))]; tensor encoder_layers_22_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_22_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(681195136)))]; tensor input_315_cast_fp16 = layer_norm(axes = input_315_axes_0, beta = encoder_layers_22_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_22_layer_norm2_weight_to_fp16, x = input_313_cast_fp16)[name = tensor("input_315_cast_fp16")]; tensor encoder_layers_22_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_22_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(681197504)))]; tensor encoder_layers_22_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_22_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(691113984)))]; tensor linear_136_cast_fp16 = linear(bias = encoder_layers_22_mlp_fc1_bias_to_fp16, weight = encoder_layers_22_mlp_fc1_weight_to_fp16, x = input_315_cast_fp16)[name = tensor("linear_136_cast_fp16")]; tensor input_319_mode_0 = const()[name = tensor("input_319_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_319_cast_fp16 = gelu(mode = input_319_mode_0, x = linear_136_cast_fp16)[name = tensor("input_319_cast_fp16")]; tensor encoder_layers_22_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_22_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(691122688)))]; tensor encoder_layers_22_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_22_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(701039168)))]; tensor linear_137_cast_fp16 = linear(bias = encoder_layers_22_mlp_fc2_bias_to_fp16, weight = encoder_layers_22_mlp_fc2_weight_to_fp16, x = input_319_cast_fp16)[name = tensor("linear_137_cast_fp16")]; tensor input_321_cast_fp16 = add(x = input_313_cast_fp16, y = linear_137_cast_fp16)[name = tensor("input_321_cast_fp16")]; tensor hidden_states_139_axes_0 = const()[name = tensor("hidden_states_139_axes_0"), val = tensor([-1])]; tensor encoder_layers_23_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_23_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(701041536)))]; tensor encoder_layers_23_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_23_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(701043904)))]; tensor hidden_states_139_cast_fp16 = layer_norm(axes = hidden_states_139_axes_0, beta = encoder_layers_23_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_23_layer_norm1_weight_to_fp16, x = input_321_cast_fp16)[name = tensor("hidden_states_139_cast_fp16")]; tensor encoder_layers_23_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_23_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(701046272)))]; tensor encoder_layers_23_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_23_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(703700544)))]; tensor linear_138_cast_fp16 = linear(bias = encoder_layers_23_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_23_self_attn_q_proj_weight_to_fp16, x = hidden_states_139_cast_fp16)[name = tensor("linear_138_cast_fp16")]; tensor encoder_layers_23_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_23_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(703702912)))]; tensor encoder_layers_23_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_23_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(706357184)))]; tensor linear_139_cast_fp16 = linear(bias = encoder_layers_23_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_23_self_attn_k_proj_weight_to_fp16, x = hidden_states_139_cast_fp16)[name = tensor("linear_139_cast_fp16")]; tensor encoder_layers_23_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_23_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(706359552)))]; tensor encoder_layers_23_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_23_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(709013824)))]; tensor linear_140_cast_fp16 = linear(bias = encoder_layers_23_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_23_self_attn_v_proj_weight_to_fp16, x = hidden_states_139_cast_fp16)[name = tensor("linear_140_cast_fp16")]; tensor var_1522 = const()[name = tensor("op_1522"), val = tensor([1, 1024, 16, 72])]; tensor var_1523_cast_fp16 = reshape(shape = var_1522, x = linear_138_cast_fp16)[name = tensor("op_1523_cast_fp16")]; tensor var_1525 = const()[name = tensor("op_1525"), val = tensor([1, 1024, 16, 72])]; tensor var_1526_cast_fp16 = reshape(shape = var_1525, x = linear_139_cast_fp16)[name = tensor("op_1526_cast_fp16")]; tensor var_1528 = const()[name = tensor("op_1528"), val = tensor([1, 1024, 16, 72])]; tensor var_1529_cast_fp16 = reshape(shape = var_1528, x = linear_140_cast_fp16)[name = tensor("op_1529_cast_fp16")]; tensor value_states_95_perm_0 = const()[name = tensor("value_states_95_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_1532_transpose_x_0 = const()[name = tensor("op_1532_transpose_x_0"), val = tensor(false)]; tensor var_1532_transpose_y_0 = const()[name = tensor("op_1532_transpose_y_0"), val = tensor(false)]; tensor transpose_127_perm_0 = const()[name = tensor("transpose_127_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_128_perm_0 = const()[name = tensor("transpose_128_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_128 = transpose(perm = transpose_128_perm_0, x = var_1526_cast_fp16)[name = tensor("transpose_148")]; tensor transpose_127 = transpose(perm = transpose_127_perm_0, x = var_1523_cast_fp16)[name = tensor("transpose_149")]; tensor var_1532_cast_fp16 = matmul(transpose_x = var_1532_transpose_x_0, transpose_y = var_1532_transpose_y_0, x = transpose_127, y = transpose_128)[name = tensor("op_1532_cast_fp16")]; tensor var_1533_to_fp16 = const()[name = tensor("op_1533_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_93_cast_fp16 = mul(x = var_1532_cast_fp16, y = var_1533_to_fp16)[name = tensor("attn_weights_93_cast_fp16")]; tensor var_1535_cast_fp16 = softmax(axis = var_11, x = attn_weights_93_cast_fp16)[name = tensor("op_1535_cast_fp16")]; tensor attn_output_93_transpose_x_0 = const()[name = tensor("attn_output_93_transpose_x_0"), val = tensor(false)]; tensor attn_output_93_transpose_y_0 = const()[name = tensor("attn_output_93_transpose_y_0"), val = tensor(false)]; tensor value_states_95_cast_fp16 = transpose(perm = value_states_95_perm_0, x = var_1529_cast_fp16)[name = tensor("transpose_150")]; tensor attn_output_93_cast_fp16 = matmul(transpose_x = attn_output_93_transpose_x_0, transpose_y = attn_output_93_transpose_y_0, x = var_1535_cast_fp16, y = value_states_95_cast_fp16)[name = tensor("attn_output_93_cast_fp16")]; tensor var_1539_perm_0 = const()[name = tensor("op_1539_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_1541 = const()[name = tensor("op_1541"), val = tensor([1, 1024, 1152])]; tensor var_1539_cast_fp16 = transpose(perm = var_1539_perm_0, x = attn_output_93_cast_fp16)[name = tensor("transpose_147")]; tensor input_325_cast_fp16 = reshape(shape = var_1541, x = var_1539_cast_fp16)[name = tensor("input_325_cast_fp16")]; tensor encoder_layers_23_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_23_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(709016192)))]; tensor encoder_layers_23_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_23_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(711670464)))]; tensor linear_141_cast_fp16 = linear(bias = encoder_layers_23_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_23_self_attn_out_proj_weight_to_fp16, x = input_325_cast_fp16)[name = tensor("linear_141_cast_fp16")]; tensor input_327_cast_fp16 = add(x = input_321_cast_fp16, y = linear_141_cast_fp16)[name = tensor("input_327_cast_fp16")]; tensor input_329_axes_0 = const()[name = tensor("input_329_axes_0"), val = tensor([-1])]; tensor encoder_layers_23_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_23_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(711672832)))]; tensor encoder_layers_23_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_23_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(711675200)))]; tensor input_329_cast_fp16 = layer_norm(axes = input_329_axes_0, beta = encoder_layers_23_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_23_layer_norm2_weight_to_fp16, x = input_327_cast_fp16)[name = tensor("input_329_cast_fp16")]; tensor encoder_layers_23_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_23_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(711677568)))]; tensor encoder_layers_23_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_23_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(721594048)))]; tensor linear_142_cast_fp16 = linear(bias = encoder_layers_23_mlp_fc1_bias_to_fp16, weight = encoder_layers_23_mlp_fc1_weight_to_fp16, x = input_329_cast_fp16)[name = tensor("linear_142_cast_fp16")]; tensor input_333_mode_0 = const()[name = tensor("input_333_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_333_cast_fp16 = gelu(mode = input_333_mode_0, x = linear_142_cast_fp16)[name = tensor("input_333_cast_fp16")]; tensor encoder_layers_23_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_23_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(721602752)))]; tensor encoder_layers_23_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_23_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(731519232)))]; tensor linear_143_cast_fp16 = linear(bias = encoder_layers_23_mlp_fc2_bias_to_fp16, weight = encoder_layers_23_mlp_fc2_weight_to_fp16, x = input_333_cast_fp16)[name = tensor("linear_143_cast_fp16")]; tensor input_335_cast_fp16 = add(x = input_327_cast_fp16, y = linear_143_cast_fp16)[name = tensor("input_335_cast_fp16")]; tensor hidden_states_145_axes_0 = const()[name = tensor("hidden_states_145_axes_0"), val = tensor([-1])]; tensor encoder_layers_24_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_24_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(731521600)))]; tensor encoder_layers_24_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_24_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(731523968)))]; tensor hidden_states_145_cast_fp16 = layer_norm(axes = hidden_states_145_axes_0, beta = encoder_layers_24_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_24_layer_norm1_weight_to_fp16, x = input_335_cast_fp16)[name = tensor("hidden_states_145_cast_fp16")]; tensor encoder_layers_24_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_24_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(731526336)))]; tensor encoder_layers_24_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_24_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(734180608)))]; tensor linear_144_cast_fp16 = linear(bias = encoder_layers_24_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_24_self_attn_q_proj_weight_to_fp16, x = hidden_states_145_cast_fp16)[name = tensor("linear_144_cast_fp16")]; tensor encoder_layers_24_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_24_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(734182976)))]; tensor encoder_layers_24_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_24_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(736837248)))]; tensor linear_145_cast_fp16 = linear(bias = encoder_layers_24_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_24_self_attn_k_proj_weight_to_fp16, x = hidden_states_145_cast_fp16)[name = tensor("linear_145_cast_fp16")]; tensor encoder_layers_24_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_24_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(736839616)))]; tensor encoder_layers_24_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_24_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(739493888)))]; tensor linear_146_cast_fp16 = linear(bias = encoder_layers_24_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_24_self_attn_v_proj_weight_to_fp16, x = hidden_states_145_cast_fp16)[name = tensor("linear_146_cast_fp16")]; tensor var_1584 = const()[name = tensor("op_1584"), val = tensor([1, 1024, 16, 72])]; tensor var_1585_cast_fp16 = reshape(shape = var_1584, x = linear_144_cast_fp16)[name = tensor("op_1585_cast_fp16")]; tensor var_1587 = const()[name = tensor("op_1587"), val = tensor([1, 1024, 16, 72])]; tensor var_1588_cast_fp16 = reshape(shape = var_1587, x = linear_145_cast_fp16)[name = tensor("op_1588_cast_fp16")]; tensor var_1590 = const()[name = tensor("op_1590"), val = tensor([1, 1024, 16, 72])]; tensor var_1591_cast_fp16 = reshape(shape = var_1590, x = linear_146_cast_fp16)[name = tensor("op_1591_cast_fp16")]; tensor value_states_99_perm_0 = const()[name = tensor("value_states_99_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_1594_transpose_x_0 = const()[name = tensor("op_1594_transpose_x_0"), val = tensor(false)]; tensor var_1594_transpose_y_0 = const()[name = tensor("op_1594_transpose_y_0"), val = tensor(false)]; tensor transpose_129_perm_0 = const()[name = tensor("transpose_129_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_130_perm_0 = const()[name = tensor("transpose_130_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_130 = transpose(perm = transpose_130_perm_0, x = var_1588_cast_fp16)[name = tensor("transpose_144")]; tensor transpose_129 = transpose(perm = transpose_129_perm_0, x = var_1585_cast_fp16)[name = tensor("transpose_145")]; tensor var_1594_cast_fp16 = matmul(transpose_x = var_1594_transpose_x_0, transpose_y = var_1594_transpose_y_0, x = transpose_129, y = transpose_130)[name = tensor("op_1594_cast_fp16")]; tensor var_1595_to_fp16 = const()[name = tensor("op_1595_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_97_cast_fp16 = mul(x = var_1594_cast_fp16, y = var_1595_to_fp16)[name = tensor("attn_weights_97_cast_fp16")]; tensor var_1597_cast_fp16 = softmax(axis = var_11, x = attn_weights_97_cast_fp16)[name = tensor("op_1597_cast_fp16")]; tensor attn_output_97_transpose_x_0 = const()[name = tensor("attn_output_97_transpose_x_0"), val = tensor(false)]; tensor attn_output_97_transpose_y_0 = const()[name = tensor("attn_output_97_transpose_y_0"), val = tensor(false)]; tensor value_states_99_cast_fp16 = transpose(perm = value_states_99_perm_0, x = var_1591_cast_fp16)[name = tensor("transpose_146")]; tensor attn_output_97_cast_fp16 = matmul(transpose_x = attn_output_97_transpose_x_0, transpose_y = attn_output_97_transpose_y_0, x = var_1597_cast_fp16, y = value_states_99_cast_fp16)[name = tensor("attn_output_97_cast_fp16")]; tensor var_1601_perm_0 = const()[name = tensor("op_1601_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_1603 = const()[name = tensor("op_1603"), val = tensor([1, 1024, 1152])]; tensor var_1601_cast_fp16 = transpose(perm = var_1601_perm_0, x = attn_output_97_cast_fp16)[name = tensor("transpose_143")]; tensor input_339_cast_fp16 = reshape(shape = var_1603, x = var_1601_cast_fp16)[name = tensor("input_339_cast_fp16")]; tensor encoder_layers_24_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_24_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(739496256)))]; tensor encoder_layers_24_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_24_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(742150528)))]; tensor linear_147_cast_fp16 = linear(bias = encoder_layers_24_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_24_self_attn_out_proj_weight_to_fp16, x = input_339_cast_fp16)[name = tensor("linear_147_cast_fp16")]; tensor input_341_cast_fp16 = add(x = input_335_cast_fp16, y = linear_147_cast_fp16)[name = tensor("input_341_cast_fp16")]; tensor input_343_axes_0 = const()[name = tensor("input_343_axes_0"), val = tensor([-1])]; tensor encoder_layers_24_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_24_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(742152896)))]; tensor encoder_layers_24_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_24_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(742155264)))]; tensor input_343_cast_fp16 = layer_norm(axes = input_343_axes_0, beta = encoder_layers_24_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_24_layer_norm2_weight_to_fp16, x = input_341_cast_fp16)[name = tensor("input_343_cast_fp16")]; tensor encoder_layers_24_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_24_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(742157632)))]; tensor encoder_layers_24_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_24_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(752074112)))]; tensor linear_148_cast_fp16 = linear(bias = encoder_layers_24_mlp_fc1_bias_to_fp16, weight = encoder_layers_24_mlp_fc1_weight_to_fp16, x = input_343_cast_fp16)[name = tensor("linear_148_cast_fp16")]; tensor input_347_mode_0 = const()[name = tensor("input_347_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_347_cast_fp16 = gelu(mode = input_347_mode_0, x = linear_148_cast_fp16)[name = tensor("input_347_cast_fp16")]; tensor encoder_layers_24_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_24_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(752082816)))]; tensor encoder_layers_24_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_24_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(761999296)))]; tensor linear_149_cast_fp16 = linear(bias = encoder_layers_24_mlp_fc2_bias_to_fp16, weight = encoder_layers_24_mlp_fc2_weight_to_fp16, x = input_347_cast_fp16)[name = tensor("linear_149_cast_fp16")]; tensor input_349_cast_fp16 = add(x = input_341_cast_fp16, y = linear_149_cast_fp16)[name = tensor("input_349_cast_fp16")]; tensor hidden_states_151_axes_0 = const()[name = tensor("hidden_states_151_axes_0"), val = tensor([-1])]; tensor encoder_layers_25_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_25_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(762001664)))]; tensor encoder_layers_25_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_25_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(762004032)))]; tensor hidden_states_151_cast_fp16 = layer_norm(axes = hidden_states_151_axes_0, beta = encoder_layers_25_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_25_layer_norm1_weight_to_fp16, x = input_349_cast_fp16)[name = tensor("hidden_states_151_cast_fp16")]; tensor encoder_layers_25_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_25_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(762006400)))]; tensor encoder_layers_25_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_25_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(764660672)))]; tensor linear_150_cast_fp16 = linear(bias = encoder_layers_25_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_25_self_attn_q_proj_weight_to_fp16, x = hidden_states_151_cast_fp16)[name = tensor("linear_150_cast_fp16")]; tensor encoder_layers_25_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_25_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(764663040)))]; tensor encoder_layers_25_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_25_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(767317312)))]; tensor linear_151_cast_fp16 = linear(bias = encoder_layers_25_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_25_self_attn_k_proj_weight_to_fp16, x = hidden_states_151_cast_fp16)[name = tensor("linear_151_cast_fp16")]; tensor encoder_layers_25_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_25_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(767319680)))]; tensor encoder_layers_25_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_25_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(769973952)))]; tensor linear_152_cast_fp16 = linear(bias = encoder_layers_25_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_25_self_attn_v_proj_weight_to_fp16, x = hidden_states_151_cast_fp16)[name = tensor("linear_152_cast_fp16")]; tensor var_1646 = const()[name = tensor("op_1646"), val = tensor([1, 1024, 16, 72])]; tensor var_1647_cast_fp16 = reshape(shape = var_1646, x = linear_150_cast_fp16)[name = tensor("op_1647_cast_fp16")]; tensor var_1649 = const()[name = tensor("op_1649"), val = tensor([1, 1024, 16, 72])]; tensor var_1650_cast_fp16 = reshape(shape = var_1649, x = linear_151_cast_fp16)[name = tensor("op_1650_cast_fp16")]; tensor var_1652 = const()[name = tensor("op_1652"), val = tensor([1, 1024, 16, 72])]; tensor var_1653_cast_fp16 = reshape(shape = var_1652, x = linear_152_cast_fp16)[name = tensor("op_1653_cast_fp16")]; tensor value_states_103_perm_0 = const()[name = tensor("value_states_103_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_1656_transpose_x_0 = const()[name = tensor("op_1656_transpose_x_0"), val = tensor(false)]; tensor var_1656_transpose_y_0 = const()[name = tensor("op_1656_transpose_y_0"), val = tensor(false)]; tensor transpose_131_perm_0 = const()[name = tensor("transpose_131_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_132_perm_0 = const()[name = tensor("transpose_132_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_132 = transpose(perm = transpose_132_perm_0, x = var_1650_cast_fp16)[name = tensor("transpose_140")]; tensor transpose_131 = transpose(perm = transpose_131_perm_0, x = var_1647_cast_fp16)[name = tensor("transpose_141")]; tensor var_1656_cast_fp16 = matmul(transpose_x = var_1656_transpose_x_0, transpose_y = var_1656_transpose_y_0, x = transpose_131, y = transpose_132)[name = tensor("op_1656_cast_fp16")]; tensor var_1657_to_fp16 = const()[name = tensor("op_1657_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_101_cast_fp16 = mul(x = var_1656_cast_fp16, y = var_1657_to_fp16)[name = tensor("attn_weights_101_cast_fp16")]; tensor var_1659_cast_fp16 = softmax(axis = var_11, x = attn_weights_101_cast_fp16)[name = tensor("op_1659_cast_fp16")]; tensor attn_output_101_transpose_x_0 = const()[name = tensor("attn_output_101_transpose_x_0"), val = tensor(false)]; tensor attn_output_101_transpose_y_0 = const()[name = tensor("attn_output_101_transpose_y_0"), val = tensor(false)]; tensor value_states_103_cast_fp16 = transpose(perm = value_states_103_perm_0, x = var_1653_cast_fp16)[name = tensor("transpose_142")]; tensor attn_output_101_cast_fp16 = matmul(transpose_x = attn_output_101_transpose_x_0, transpose_y = attn_output_101_transpose_y_0, x = var_1659_cast_fp16, y = value_states_103_cast_fp16)[name = tensor("attn_output_101_cast_fp16")]; tensor var_1663_perm_0 = const()[name = tensor("op_1663_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_1665 = const()[name = tensor("op_1665"), val = tensor([1, 1024, 1152])]; tensor var_1663_cast_fp16 = transpose(perm = var_1663_perm_0, x = attn_output_101_cast_fp16)[name = tensor("transpose_139")]; tensor input_353_cast_fp16 = reshape(shape = var_1665, x = var_1663_cast_fp16)[name = tensor("input_353_cast_fp16")]; tensor encoder_layers_25_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_25_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(769976320)))]; tensor encoder_layers_25_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_25_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(772630592)))]; tensor linear_153_cast_fp16 = linear(bias = encoder_layers_25_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_25_self_attn_out_proj_weight_to_fp16, x = input_353_cast_fp16)[name = tensor("linear_153_cast_fp16")]; tensor input_355_cast_fp16 = add(x = input_349_cast_fp16, y = linear_153_cast_fp16)[name = tensor("input_355_cast_fp16")]; tensor input_357_axes_0 = const()[name = tensor("input_357_axes_0"), val = tensor([-1])]; tensor encoder_layers_25_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_25_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(772632960)))]; tensor encoder_layers_25_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_25_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(772635328)))]; tensor input_357_cast_fp16 = layer_norm(axes = input_357_axes_0, beta = encoder_layers_25_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_25_layer_norm2_weight_to_fp16, x = input_355_cast_fp16)[name = tensor("input_357_cast_fp16")]; tensor encoder_layers_25_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_25_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(772637696)))]; tensor encoder_layers_25_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_25_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(782554176)))]; tensor linear_154_cast_fp16 = linear(bias = encoder_layers_25_mlp_fc1_bias_to_fp16, weight = encoder_layers_25_mlp_fc1_weight_to_fp16, x = input_357_cast_fp16)[name = tensor("linear_154_cast_fp16")]; tensor input_361_mode_0 = const()[name = tensor("input_361_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_361_cast_fp16 = gelu(mode = input_361_mode_0, x = linear_154_cast_fp16)[name = tensor("input_361_cast_fp16")]; tensor encoder_layers_25_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_25_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(782562880)))]; tensor encoder_layers_25_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_25_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(792479360)))]; tensor linear_155_cast_fp16 = linear(bias = encoder_layers_25_mlp_fc2_bias_to_fp16, weight = encoder_layers_25_mlp_fc2_weight_to_fp16, x = input_361_cast_fp16)[name = tensor("linear_155_cast_fp16")]; tensor input_363_cast_fp16 = add(x = input_355_cast_fp16, y = linear_155_cast_fp16)[name = tensor("input_363_cast_fp16")]; tensor hidden_states_157_axes_0 = const()[name = tensor("hidden_states_157_axes_0"), val = tensor([-1])]; tensor encoder_layers_26_layer_norm1_weight_to_fp16 = const()[name = tensor("encoder_layers_26_layer_norm1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(792481728)))]; tensor encoder_layers_26_layer_norm1_bias_to_fp16 = const()[name = tensor("encoder_layers_26_layer_norm1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(792484096)))]; tensor hidden_states_157_cast_fp16 = layer_norm(axes = hidden_states_157_axes_0, beta = encoder_layers_26_layer_norm1_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_26_layer_norm1_weight_to_fp16, x = input_363_cast_fp16)[name = tensor("hidden_states_157_cast_fp16")]; tensor encoder_layers_26_self_attn_q_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_26_self_attn_q_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(792486464)))]; tensor encoder_layers_26_self_attn_q_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_26_self_attn_q_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(795140736)))]; tensor linear_156_cast_fp16 = linear(bias = encoder_layers_26_self_attn_q_proj_bias_to_fp16, weight = encoder_layers_26_self_attn_q_proj_weight_to_fp16, x = hidden_states_157_cast_fp16)[name = tensor("linear_156_cast_fp16")]; tensor encoder_layers_26_self_attn_k_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_26_self_attn_k_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(795143104)))]; tensor encoder_layers_26_self_attn_k_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_26_self_attn_k_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(797797376)))]; tensor linear_157_cast_fp16 = linear(bias = encoder_layers_26_self_attn_k_proj_bias_to_fp16, weight = encoder_layers_26_self_attn_k_proj_weight_to_fp16, x = hidden_states_157_cast_fp16)[name = tensor("linear_157_cast_fp16")]; tensor encoder_layers_26_self_attn_v_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_26_self_attn_v_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(797799744)))]; tensor encoder_layers_26_self_attn_v_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_26_self_attn_v_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(800454016)))]; tensor linear_158_cast_fp16 = linear(bias = encoder_layers_26_self_attn_v_proj_bias_to_fp16, weight = encoder_layers_26_self_attn_v_proj_weight_to_fp16, x = hidden_states_157_cast_fp16)[name = tensor("linear_158_cast_fp16")]; tensor var_1708 = const()[name = tensor("op_1708"), val = tensor([1, 1024, 16, 72])]; tensor var_1709_cast_fp16 = reshape(shape = var_1708, x = linear_156_cast_fp16)[name = tensor("op_1709_cast_fp16")]; tensor var_1711 = const()[name = tensor("op_1711"), val = tensor([1, 1024, 16, 72])]; tensor var_1712_cast_fp16 = reshape(shape = var_1711, x = linear_157_cast_fp16)[name = tensor("op_1712_cast_fp16")]; tensor var_1714 = const()[name = tensor("op_1714"), val = tensor([1, 1024, 16, 72])]; tensor var_1715_cast_fp16 = reshape(shape = var_1714, x = linear_158_cast_fp16)[name = tensor("op_1715_cast_fp16")]; tensor value_states_perm_0 = const()[name = tensor("value_states_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_1718_transpose_x_0 = const()[name = tensor("op_1718_transpose_x_0"), val = tensor(false)]; tensor var_1718_transpose_y_0 = const()[name = tensor("op_1718_transpose_y_0"), val = tensor(false)]; tensor transpose_133_perm_0 = const()[name = tensor("transpose_133_perm_0"), val = tensor([0, 2, -3, -1])]; tensor transpose_134_perm_0 = const()[name = tensor("transpose_134_perm_0"), val = tensor([0, 2, -1, -3])]; tensor transpose_134 = transpose(perm = transpose_134_perm_0, x = var_1712_cast_fp16)[name = tensor("transpose_136")]; tensor transpose_133 = transpose(perm = transpose_133_perm_0, x = var_1709_cast_fp16)[name = tensor("transpose_137")]; tensor var_1718_cast_fp16 = matmul(transpose_x = var_1718_transpose_x_0, transpose_y = var_1718_transpose_y_0, x = transpose_133, y = transpose_134)[name = tensor("op_1718_cast_fp16")]; tensor var_1719_to_fp16 = const()[name = tensor("op_1719_to_fp16"), val = tensor(0x1.e2cp-4)]; tensor attn_weights_105_cast_fp16 = mul(x = var_1718_cast_fp16, y = var_1719_to_fp16)[name = tensor("attn_weights_105_cast_fp16")]; tensor var_1721_cast_fp16 = softmax(axis = var_11, x = attn_weights_105_cast_fp16)[name = tensor("op_1721_cast_fp16")]; tensor attn_output_105_transpose_x_0 = const()[name = tensor("attn_output_105_transpose_x_0"), val = tensor(false)]; tensor attn_output_105_transpose_y_0 = const()[name = tensor("attn_output_105_transpose_y_0"), val = tensor(false)]; tensor value_states_cast_fp16 = transpose(perm = value_states_perm_0, x = var_1715_cast_fp16)[name = tensor("transpose_138")]; tensor attn_output_105_cast_fp16 = matmul(transpose_x = attn_output_105_transpose_x_0, transpose_y = attn_output_105_transpose_y_0, x = var_1721_cast_fp16, y = value_states_cast_fp16)[name = tensor("attn_output_105_cast_fp16")]; tensor var_1725_perm_0 = const()[name = tensor("op_1725_perm_0"), val = tensor([0, 2, 1, 3])]; tensor var_1727 = const()[name = tensor("op_1727"), val = tensor([1, 1024, 1152])]; tensor var_1725_cast_fp16 = transpose(perm = var_1725_perm_0, x = attn_output_105_cast_fp16)[name = tensor("transpose_135")]; tensor input_367_cast_fp16 = reshape(shape = var_1727, x = var_1725_cast_fp16)[name = tensor("input_367_cast_fp16")]; tensor encoder_layers_26_self_attn_out_proj_weight_to_fp16 = const()[name = tensor("encoder_layers_26_self_attn_out_proj_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(800456384)))]; tensor encoder_layers_26_self_attn_out_proj_bias_to_fp16 = const()[name = tensor("encoder_layers_26_self_attn_out_proj_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(803110656)))]; tensor linear_159_cast_fp16 = linear(bias = encoder_layers_26_self_attn_out_proj_bias_to_fp16, weight = encoder_layers_26_self_attn_out_proj_weight_to_fp16, x = input_367_cast_fp16)[name = tensor("linear_159_cast_fp16")]; tensor input_369_cast_fp16 = add(x = input_363_cast_fp16, y = linear_159_cast_fp16)[name = tensor("input_369_cast_fp16")]; tensor input_371_axes_0 = const()[name = tensor("input_371_axes_0"), val = tensor([-1])]; tensor encoder_layers_26_layer_norm2_weight_to_fp16 = const()[name = tensor("encoder_layers_26_layer_norm2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(803113024)))]; tensor encoder_layers_26_layer_norm2_bias_to_fp16 = const()[name = tensor("encoder_layers_26_layer_norm2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(803115392)))]; tensor input_371_cast_fp16 = layer_norm(axes = input_371_axes_0, beta = encoder_layers_26_layer_norm2_bias_to_fp16, epsilon = var_5_to_fp16, gamma = encoder_layers_26_layer_norm2_weight_to_fp16, x = input_369_cast_fp16)[name = tensor("input_371_cast_fp16")]; tensor encoder_layers_26_mlp_fc1_weight_to_fp16 = const()[name = tensor("encoder_layers_26_mlp_fc1_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(803117760)))]; tensor encoder_layers_26_mlp_fc1_bias_to_fp16 = const()[name = tensor("encoder_layers_26_mlp_fc1_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(813034240)))]; tensor linear_160_cast_fp16 = linear(bias = encoder_layers_26_mlp_fc1_bias_to_fp16, weight = encoder_layers_26_mlp_fc1_weight_to_fp16, x = input_371_cast_fp16)[name = tensor("linear_160_cast_fp16")]; tensor input_375_mode_0 = const()[name = tensor("input_375_mode_0"), val = tensor("TANH_APPROXIMATION")]; tensor input_375_cast_fp16 = gelu(mode = input_375_mode_0, x = linear_160_cast_fp16)[name = tensor("input_375_cast_fp16")]; tensor encoder_layers_26_mlp_fc2_weight_to_fp16 = const()[name = tensor("encoder_layers_26_mlp_fc2_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(813042944)))]; tensor encoder_layers_26_mlp_fc2_bias_to_fp16 = const()[name = tensor("encoder_layers_26_mlp_fc2_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(822959424)))]; tensor linear_161_cast_fp16 = linear(bias = encoder_layers_26_mlp_fc2_bias_to_fp16, weight = encoder_layers_26_mlp_fc2_weight_to_fp16, x = input_375_cast_fp16)[name = tensor("linear_161_cast_fp16")]; tensor input_cast_fp16 = add(x = input_369_cast_fp16, y = linear_161_cast_fp16)[name = tensor("input_cast_fp16")]; tensor var_1753_axes_0 = const()[name = tensor("op_1753_axes_0"), val = tensor([-1])]; tensor post_layernorm_weight_to_fp16 = const()[name = tensor("post_layernorm_weight_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(822961792)))]; tensor post_layernorm_bias_to_fp16 = const()[name = tensor("post_layernorm_bias_to_fp16"), val = tensor(BLOBFILE(path = tensor("@model_path/weights/weight.bin"), offset = tensor(822964160)))]; tensor var_1748_to_fp16 = const()[name = tensor("op_1748_to_fp16"), val = tensor(0x1.1p-20)]; tensor var_1753_cast_fp16 = layer_norm(axes = var_1753_axes_0, beta = post_layernorm_bias_to_fp16, epsilon = var_1748_to_fp16, gamma = post_layernorm_weight_to_fp16, x = input_cast_fp16)[name = tensor("op_1753_cast_fp16")]; tensor var_1753_cast_fp16_to_fp32_dtype_0 = const()[name = tensor("op_1753_cast_fp16_to_fp32_dtype_0"), val = tensor("fp32")]; tensor output = cast(dtype = var_1753_cast_fp16_to_fp32_dtype_0, x = var_1753_cast_fp16)[name = tensor("cast_135")]; } -> (output); }