@misc{indiciae3332c9b7c942, title = {Less Data, Faster Training: repeating smaller datasets speeds up learning via sampling biases}, author = {Jingwen Liu and Ezra Edelman and Surbhi Goel and Bingbin Liu}, year = {2026}, url = {https://arxiv.org/abs/2605.20314}, note = {Source identifier: 2605.20314} }