@misc{indiciae95d0bddca376, title = {Why Does Adaptive Batching Help LLM Pretraining? A Perspective from Unbounded Variance}, author = {Arda Fazla and Antesh Upadhyay and Ege C. Kaya and M. Berk Sahin and Abolfazl Hashemi}, year = {2026}, url = {https://arxiv.org/abs/2610.02355}, note = {Source identifier: 2610.02355} }