LangXAI is a framework that integrates Explainable Artificial Intelligence (XAI) with advanced vision models to generate textual explanations for visual recognition tasks. Despite XAI advancements, an understanding gap persists for end-users with limited domain knowledge in artificial intelligence and computer vision. LangXAI addresses this by furnishing text-based explanations for classification, object detection, and semantic segmentation model outputs to end-users. Preliminary results demonstrate LangXAI's enhanced plausibility, with high BERTScore across tasks, fostering a more transparent and reliable AI framework on vision tasks for end-users.
@article{arxiv.2402.12525,
title = {LangXAI: Integrating Large Vision Models for Generating Textual Explanations to Enhance Explainability in Visual Perception Tasks},
author = {Truong Thanh Hung Nguyen and Tobias Clement and Phuc Truong Loc Nguyen and Nils Kemmerzell and Van Binh Truong and Vo Thanh Khang Nguyen and Mohamed Abdelaal and Hung Cao},
journal= {arXiv preprint arXiv:2402.12525},
year = {2024}
}