Make NPROC_PER_NODE customizable in run1000.sh and speedrun.sh

2025-12-06 04:12:13 +00:00 · 2025-11-04 22:16:08 -08:00 · 2025-11-04 22:16:08 -08:00 · a2d61393ee
commit a2d61393ee
parent 885a4f25e7
2 changed files with 2 additions and 2 deletions
--- a/run1000.sh
+++ b/run1000.sh
@ -72,7 +72,7 @@ python -m scripts.tok_eval
 # 5) That's it, everything else (e.g. the learning rates) is adjusted automatically by the training script.

 # Number of processes/GPUs to use
-NPROC_PER_NODE=8
+NPROC_PER_NODE=${NPROC_PER_NODE:-8}

 torchrun --standalone --nproc_per_node=$NPROC_PER_NODE -m scripts.base_train -- --depth=32 --device_batch_size=8 --run=$WANDB_RUN
 torchrun --standalone --nproc_per_node=$NPROC_PER_NODE -m scripts.base_loss
--- a/speedrun.sh
+++ b/speedrun.sh
@ -83,7 +83,7 @@ echo "Waiting for dataset download to complete..."
 wait $DATASET_DOWNLOAD_PID

 # Number of processes/GPUs to use
-NPROC_PER_NODE=8
+NPROC_PER_NODE=${NPROC_PER_NODE:-8}

 # pretrain the d20 model
 torchrun --standalone --nproc_per_node=$NPROC_PER_NODE -m scripts.base_train -- --depth=20 --run=$WANDB_RUN