mirror of
https://github.com/NVIDIA/TensorRT-LLM.git
synced 2026-01-14 06:27:45 +08:00
Signed-off-by: Gal Hubara Agam <96368689+galagam@users.noreply.github.com> Signed-off-by: Neta Zmora <96238833+nzmora-nvidia@users.noreply.github.com> Signed-off-by: Lucas Liebenwein <11156568+lucaslie@users.noreply.github.com> Signed-off-by: nvchenghaoz <211069071+nvchenghaoz@users.noreply.github.com> Signed-off-by: Frida Hou <201670829+Fridah-nv@users.noreply.github.com> Signed-off-by: greg-kwasniewski1 <213329731+greg-kwasniewski1@users.noreply.github.com> Signed-off-by: Suyog Gupta <41447211+suyoggupta@users.noreply.github.com> Co-authored-by: Gal Hubara-Agam <96368689+galagam@users.noreply.github.com> Co-authored-by: Neta Zmora <nzmora@nvidia.com> Co-authored-by: nvchenghaoz <211069071+nvchenghaoz@users.noreply.github.com> Co-authored-by: Frida Hou <201670829+Fridah-nv@users.noreply.github.com> Co-authored-by: Suyog Gupta <41447211+suyoggupta@users.noreply.github.com> Co-authored-by: Grzegorz Kwasniewski <213329731+greg-kwasniewski1@users.noreply.github.com>
41 lines
1.4 KiB
JSON
41 lines
1.4 KiB
JSON
{
|
|
"version": "0.2.0",
|
|
"configurations": [
|
|
{
|
|
"name": "build_and_run_ad.py",
|
|
"type": "debugpy",
|
|
"request": "launch",
|
|
"program": "build_and_run_ad.py",
|
|
"args": [
|
|
"--model=meta-llama/Meta-Llama-3.1-8B-Instruct",
|
|
"--args.world-size=2",
|
|
"--args.runtime=demollm",
|
|
"--args.compile-backend=torch-simple",
|
|
"--args.attn-page-size=16",
|
|
"--args.attn-backend=flashinfer",
|
|
"--args.model-factory=AutoModelForCausalLM",
|
|
"--benchmark.enabled=false",
|
|
"--prompt.batch-size=2",
|
|
"--args.model-kwargs.num-hidden-layers=3",
|
|
"--args.model-kwargs.num-attention-heads=32",
|
|
"--prompt.sp-kwargs.max-tokens=128",
|
|
// "--dry-run", // uncomment to print the final config and return
|
|
],
|
|
"console": "integratedTerminal",
|
|
"justMyCode": false,
|
|
"cwd": "${workspaceFolder}/examples/auto_deploy"
|
|
},
|
|
{
|
|
"name": "Python: Debug Tests",
|
|
"type": "debugpy",
|
|
"request": "launch",
|
|
"program": "${file}",
|
|
"purpose": [
|
|
"debug-test",
|
|
],
|
|
"console": "integratedTerminal",
|
|
"justMyCode": false
|
|
},
|
|
]
|
|
}
|