@misc{indiciae414b60e9a21b, title = {VisLens: Single-Pass Interpretable Visual Search for Multimodal LLMs}, author = {Jingyi He and Sanghwan Kim and Zeynep Akata}, year = {2026}, url = {https://arxiv.org/abs/2608.30705}, note = {Source identifier: 2608.30705} }