@misc{indiciaec182939b92c9, title = {Prefixing Attention Sinks can Mitigate Activation Outliers for Large Language Model Quantization}, author = {Seungwoo Son and Wonpyo Park and Woohyun Han and Kyuyeun Kim and Jaeho Lee}, year = {2024}, url = {https://arxiv.org/abs/2406.12016}, note = {Source identifier: 2406.12016} }