@misc{indiciaeae9e0204094b, title = {Harnessing Multimodal Large Language Models for Training-Free Human-Object Interaction Detection}, author = {Zhaolin Cai and Huiyu Duan and Liu Yang and Yanjun Qin and Bo Ai and Wei Chen and Xiongkuo Min and Guangtao Zhai}, year = {2026}, url = {https://arxiv.org/abs/2610.06394}, note = {Source identifier: 2610.06394} }