@article{cyclicpolicydistillationsampleefficient, title = {Cyclic Policy Distillation: Sample-Efficient Sim-to-Real Reinforcement Learning with Domain Randomization}, author = {Yuki Kadokawa and Lingwei Zhu and Yoshihisa Tsurumine and Takamitsu Matsubara}, year = {2022}, eprint = {2207.14561}, archivePrefix = {arXiv}, url = {https://arxiv.org/abs/2207.14561v2}, }