onnxruntime/orttraining/tools/ci_test/run_bert_perf_test.py
Ashrit Shetty df873177eb
Update win-ort-main to tip main 250116 (#23398)
### Description
This PR is to update the win-ort-main branch to the tip main
branch as of 2025-01-16.

### Motivation and Context
This update includes the OpenVino fix for debug builds.

---------

Signed-off-by: Liqun Fu <liqfu@microsoft.com>
Signed-off-by: Liqun Fu <liqun.fu@microsoft.com>
Signed-off-by: Junze Wu <junze.wu@intel.com>
Signed-off-by: dependabot[bot] <support@github.com>
Signed-off-by: Jianhui Dai <jianhui.j.dai@intel.com>
Co-authored-by: Yueqing Zhang <yuz75@Pitt.edu>
Co-authored-by: amancini-N <63410090+amancini-N@users.noreply.github.com>
Co-authored-by: Adrian Lizarraga <adlizarraga@microsoft.com>
Co-authored-by: liqun Fu <liqfu@microsoft.com>
Co-authored-by: Guenther Schmuelling <guschmue@microsoft.com>
Co-authored-by: Yifan Li <109183385+yf711@users.noreply.github.com>
Co-authored-by: yf711 <yifanl@microsoft.com>
Co-authored-by: Wanming Lin <wanming.lin@intel.com>
Co-authored-by: wejoncy <wejoncy@163.com>
Co-authored-by: wejoncy <wejoncy@.com>
Co-authored-by: Scott McKay <skottmckay@gmail.com>
Co-authored-by: Changming Sun <chasun@microsoft.com>
Co-authored-by: Jean-Michaël Celerier <jeanmichael.celerier+github@gmail.com>
Co-authored-by: Dmitry Deshevoy <mityada@gmail.com>
Co-authored-by: xhcao <xinghua.cao@intel.com>
Co-authored-by: Yueqing Zhang <yueqingz@amd.com>
Co-authored-by: Yulong Wang <7679871+fs-eire@users.noreply.github.com>
Co-authored-by: Jiajia Qin <jiajiaqin@microsoft.com>
Co-authored-by: Wu, Junze <junze.wu@intel.com>
Co-authored-by: Jian Chen <cjian@microsoft.com>
Co-authored-by: github-actions[bot] <41898282+github-actions[bot]@users.noreply.github.com>
Co-authored-by: Matthieu Darbois <mayeut@users.noreply.github.com>
Co-authored-by: Prathik Rao <prathik.rao@gmail.com>
Co-authored-by: wonchung-microsoft <wonchung@microsoft.com>
Co-authored-by: Vincent Wang <wangwchpku@outlook.com>
Co-authored-by: PARK DongHa <luncliff@gmail.com>
Co-authored-by: Hector Li <hecli@microsoft.com>
Co-authored-by: Sam Webster <13457618+samwebster@users.noreply.github.com>
Co-authored-by: Adrian Lizarraga <adrianlm2@gmail.com>
Co-authored-by: Preetha Veeramalai <preetha.veeramalai@intel.com>
Co-authored-by: jatinwadhwa921 <jatin.wadhwa@intel.com>
Co-authored-by: Satya Kumar Jandhyala <satya.k.jandhyala@gmail.com>
Co-authored-by: Corentin Maravat <101636442+cocotdf@users.noreply.github.com>
Co-authored-by: Xiaoyu <85524621+xiaoyu-work@users.noreply.github.com>
Co-authored-by: Tianlei Wu <tlwu@microsoft.com>
Co-authored-by: dependabot[bot] <49699333+dependabot[bot]@users.noreply.github.com>
Co-authored-by: Jie Chen <jie.a.chen@intel.com>
Co-authored-by: Jianhui Dai <jianhui.j.dai@intel.com>
Co-authored-by: Copilot <175728472+Copilot@users.noreply.github.com>
Co-authored-by: Edward Chen <18449977+edgchen1@users.noreply.github.com>
Co-authored-by: Baiju Meswani <bmeswani@microsoft.com>
Co-authored-by: kunal-vaishnavi <115581922+kunal-vaishnavi@users.noreply.github.com>
Co-authored-by: Justin Chu <justinchuby@users.noreply.github.com>
Co-authored-by: Copilot Autofix powered by AI <62310815+github-advanced-security[bot]@users.noreply.github.com>
Co-authored-by: Ted Themistokleous <107195283+TedThemistokleous@users.noreply.github.com>
Co-authored-by: Jeff Daily <jeff.daily@amd.com>
Co-authored-by: Artur Wojcik <artur.wojcik@outlook.com>
Co-authored-by: Ted Themistokleous <tedthemistokleous@amd.com>
Co-authored-by: Xinya Zhang <Xinya.Zhang@amd.com>
Co-authored-by: ikalinic <ilija.kalinic@amd.com>
Co-authored-by: sstamenk <sstamenk@amd.com>
Co-authored-by: Yi-Hong Lyu <yilyu@microsoft.com>
Co-authored-by: Ti-Tai Wang <titaiwang@microsoft.com>
2025-01-16 15:20:25 -08:00

113 lines
3.7 KiB
Python

#!/usr/bin/env python3
# Copyright (c) Microsoft Corporation. All rights reserved.
# Licensed under the MIT License.
import argparse
import json
import os
import subprocess
import sys
from collections import namedtuple
SCRIPT_DIR = os.path.realpath(os.path.dirname(__file__))
def parse_args():
parser = argparse.ArgumentParser(description="Runs BERT performance tests.")
parser.add_argument("--binary_dir", required=True, help="Path to the ORT binary directory.")
parser.add_argument("--training_data_root", required=True, help="Path to the training data root directory.")
parser.add_argument("--model_root", required=True, help="Path to the model root directory.")
parser.add_argument(
"--gpu_sku",
choices=["V100_16G", "MI100_32G"],
default="V100_16G",
required=False,
help="GPU model (e.g. V100_16G, MI100_32G).",
)
return parser.parse_args()
# using the same params from "GitHub Master Merge Schedule" in OneNotes
def main():
args = parse_args()
Config = namedtuple(
"Config", ["use_mixed_precision", "max_seq_length", "batch_size", "max_predictions_per_seq", "expected_perf"]
)
configs = {}
configs["V100_16G"] = [
Config(True, 128, 76, 20, -1.0),
Config(True, 512, 11, 80, -1.0),
Config(False, 128, 39, 20, -1.0),
Config(False, 512, 6, 80, -1.0),
]
configs["MI100_32G"] = [
Config(True, 128, 128, 20, 240),
]
# run BERT training
for c in configs[args.gpu_sku]:
model = "bert-large-uncased_L_24_H_1024_A_16_V_30528_S_512_Dp_0.1_optimized_layer_norm_opset12"
precision_prefix = "fp16" if c.use_mixed_precision else "fp32"
print(
"######## testing name - "
+ ("fp16-" if c.use_mixed_precision else "fp32-")
+ str(c.max_seq_length)
+ " ##############"
)
cmds = [
os.path.join(args.binary_dir, "onnxruntime_training_bert"),
"--model_name",
os.path.join(args.model_root, f"nv/bert-large/{model}"),
"--train_data_dir",
os.path.join(args.training_data_root, str(c.max_seq_length), "books_wiki_en_corpus/train"),
"--test_data_dir",
os.path.join(args.training_data_root, str(c.max_seq_length), "books_wiki_en_corpus/test"),
"--train_batch_size",
str(c.batch_size),
"--mode",
"train",
"--max_seq_length",
str(c.max_seq_length),
"--num_train_steps",
"640",
"--display_loss_steps",
"5",
"--optimizer",
"Lamb",
"--learning_rate",
"3e-3",
"--warmup_ratio",
"0.2843",
"--warmup_mode",
"Poly",
"--gradient_accumulation_steps",
"1",
"--max_predictions_per_seq",
str(c.max_predictions_per_seq),
"--lambda",
"0",
"--use_nccl",
"--perf_output_dir",
os.path.join(SCRIPT_DIR, "results"),
]
if c.use_mixed_precision:
(cmds.append("--use_mixed_precision"),)
(cmds.append("--allreduce_in_fp16"),)
subprocess.run(cmds).check_returncode() # noqa: PLW1510
if c.expected_perf > 0.0:
json_filename = (
f"onnxruntime_perf_metrics_{model}.onnx_bert_{precision_prefix}_{c.max_seq_length}_Lamb.json"
)
with open(os.path.join(SCRIPT_DIR, "results", json_filename)) as json_file:
results = json.load(json_file)
assert results["EndToEndThroughput"] > 0.98 * c.expected_perf
return 0
if __name__ == "__main__":
sys.exit(main())