@misc{indiciaec9d292add81a, title = {Can Reinforcement Learning Unlock the Hidden Dangers in Aligned Large Language Models?}, author = {Mohammad Bahrami Karkevandi and Nishant Vishwamitra and Peyman Najafirad}, year = {2024}, url = {https://arxiv.org/abs/2408.02651}, note = {Source identifier: 2408.02651} }