@misc{indiciaeafeafca411a1, title = {Safety-Tuned LLaMAs: Lessons From Improving the Safety of Large Language Models that Follow Instructions}, author = {Federico Bianchi and Mirac Suzgun and Giuseppe Attanasio and Paul Röttger and Dan Jurafsky and Tatsunori Hashimoto and James Zou}, year = {2024}, url = {https://arxiv.org/abs/2309.07875}, note = {Source identifier: 2309.07875} }