@misc{indiciaef32e7f49e461, title = {Soft Prompt Threats: Attacking Safety Alignment and Unlearning in Open-Source LLMs through the Embedding Space}, author = {Leo Schwinn and David Dobre and Sophie Xhonneux and Gauthier Gidel and Stephan Gunnemann}, year = {2025}, url = {https://arxiv.org/abs/2402.09063}, note = {Source identifier: 2402.09063} }