@article{reasoningunder1billionmemoryaugmented, title = {Reasoning Under 1 Billion: Memory-Augmented Reinforcement Learning for Large Language Models}, author = {Hung Le and Dai Do and Dung Nguyen and Svetha Venkatesh}, year = {2025}, eprint = {2504.02273}, archivePrefix = {arXiv}, url = {https://arxiv.org/abs/2504.02273v1}, }