bitsandbytes==0.48.2
triton>=3.0.0
xformers>=0.0.23.post1
liger-kernel==0.6.3
packaging==23.2
huggingface_hub>=0.36.0
peft>=0.17.1
tokenizers>=0.22.1
transformers==4.57.1
accelerate==1.11.0
datasets==4.4.1
trl==0.25.0
hf_xet==1.2.0
kernels>=0.9.0
trackio
optimum==1.16.2
hf_transfer
sentencepiece
gradio==5.49.1
modal==1.0.2
pydantic>=2.10.6
addict
fire
PyYAML>=6.0
requests
wandb
einops
colorama
numba>=0.61.2
numpy>=2.2.6
evaluate==0.4.1
scipy
scikit-learn==1.4.2
nvidia-ml-py==12.560.30
art
tensorboard
python-dotenv==1.0.1
s3fs>=2024.5.0
gcsfs>=2025.3.0
adlfs>=2024.5.0
ocifs==1.3.2
zstandard==0.22.0
fastcore
lm_eval==0.4.7
langdetect==1.0.9
immutabledict==4.2.0
antlr4-python3-runtime==4.13.2
torchao==0.13.0
openenv-core==0.1.0
schedulefree==1.4.1
axolotl-contribs-lgpl==0.0.7
axolotl-contribs-mit==0.0.5
posthog==6.7.11
mistral-common==1.8.5
torch==2.8.0

[apollo]
apollo-torch

[auto-gptq]
auto-gptq==0.5.1

[deepspeed]
deepspeed==0.18.2
deepspeed-kernels

[fbgemm-gpu]
fbgemm-gpu-genai==1.3.0

[flash-attn]
flash-attn==2.8.3

[galore]
galore_torch

[llmcompressor]
llmcompressor==0.5.1

[mamba-ssm]
mamba-ssm==1.2.0.post1
causal_conv1d

[mlflow]
mlflow

[opentelemetry]
opentelemetry-api
opentelemetry-sdk
opentelemetry-exporter-prometheus
prometheus-client

[optimizers]
galore_torch
apollo-torch
lomo-optim==0.1.1
torch-optimi==0.2.1
came_pytorch==0.1.3

[ray]
ray[train]

[ring-flash-attn]
flash-attn==2.8.3
ring-flash-attn>=0.1.7

[vllm]
vllm==0.11.0
