@article{sampleefficientonpolicyimitationlearning, title = {Mimicking Better by Matching the Approximate Action Distribution}, author = {João A. Cândido Ramos and Lionel Blondé and Naoya Takeishi and Alexandros Kalousis}, year = {2023}, eprint = {2306.09805}, archivePrefix = {arXiv}, url = {https://arxiv.org/abs/2306.09805v3}, }