@article{rosedoesntdothatboostingthesafetyof, title = {ROSE Doesn't Do That: Boosting the Safety of Instruction-Tuned Large Language Models with Reverse Prompt Contrastive Decoding}, author = {Qihuang Zhong and Liang Ding and Juhua Liu and Bo Du and DaCheng Tao}, year = {2024}, eprint = {2402.11889}, archivePrefix = {arXiv}, url = {https://arxiv.org/abs/2402.11889v2}, }