@inproceedings{jong-koo-2026-including,
title = "Does Including Images Improve Multimodal Language Models' Accuracy on Item Discrimination Estimation",
author = "Jong, Jae Jun and
Koo, Miryeong",
editor = "Wilson, Joshua and
Ormerod, Christopher and
Beiting-Parrish, Magdalen",
booktitle = "Proceedings of the Artificial Intelligence in Measurement and Education Conference ({AIME}-Con): Full Papers",
month = oct,
year = "2026",
address = "Wyndham Grand Pittsburgh Downtown, Pittsburgh, Pennsylvania, United States",
publisher = "National Council on Measurement in Education (NCME)",
url = "https://aclanthology.org/2026.aimecon-main.71/",
pages = "631--638",
ISBN = "979-8-9983004-0-0",
abstract = "This study measures the effectiveness of including images to predict discrimination parameters of reading items. The results suggest that providing images to language models is beneficial. By using both image and text, Lasso regressions (Tibshirani, 1996) could explain data better and random forests (Breiman, 2001) showed higher accuracy."
}<?xml version="1.0" encoding="UTF-8"?>
<modsCollection xmlns="http://www.loc.gov/mods/v3">
<mods ID="jong-koo-2026-including">
<titleInfo>
<title>Does Including Images Improve Multimodal Language Models’ Accuracy on Item Discrimination Estimation</title>
</titleInfo>
<name type="personal">
<namePart type="given">Jae</namePart>
<namePart type="given">Jun</namePart>
<namePart type="family">Jong</namePart>
<role>
<roleTerm authority="marcrelator" type="text">author</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Miryeong</namePart>
<namePart type="family">Koo</namePart>
<role>
<roleTerm authority="marcrelator" type="text">author</roleTerm>
</role>
</name>
<originInfo>
<dateIssued>2026-10</dateIssued>
</originInfo>
<typeOfResource>text</typeOfResource>
<relatedItem type="host">
<titleInfo>
<title>Proceedings of the Artificial Intelligence in Measurement and Education Conference (AIME-Con): Full Papers</title>
</titleInfo>
<name type="personal">
<namePart type="given">Joshua</namePart>
<namePart type="family">Wilson</namePart>
<role>
<roleTerm authority="marcrelator" type="text">editor</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Christopher</namePart>
<namePart type="family">Ormerod</namePart>
<role>
<roleTerm authority="marcrelator" type="text">editor</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Magdalen</namePart>
<namePart type="family">Beiting-Parrish</namePart>
<role>
<roleTerm authority="marcrelator" type="text">editor</roleTerm>
</role>
</name>
<originInfo>
<publisher>National Council on Measurement in Education (NCME)</publisher>
<place>
<placeTerm type="text">Wyndham Grand Pittsburgh Downtown, Pittsburgh, Pennsylvania, United States</placeTerm>
</place>
</originInfo>
<genre authority="marcgt">conference publication</genre>
<identifier type="isbn">979-8-9983004-0-0</identifier>
</relatedItem>
<abstract>This study measures the effectiveness of including images to predict discrimination parameters of reading items. The results suggest that providing images to language models is beneficial. By using both image and text, Lasso regressions (Tibshirani, 1996) could explain data better and random forests (Breiman, 2001) showed higher accuracy.</abstract>
<identifier type="citekey">jong-koo-2026-including</identifier>
<location>
<url>https://aclanthology.org/2026.aimecon-main.71/</url>
</location>
<part>
<date>2026-10</date>
<extent unit="page">
<start>631</start>
<end>638</end>
</extent>
</part>
</mods>
</modsCollection>
%0 Conference Proceedings
%T Does Including Images Improve Multimodal Language Models’ Accuracy on Item Discrimination Estimation
%A Jong, Jae Jun
%A Koo, Miryeong
%Y Wilson, Joshua
%Y Ormerod, Christopher
%Y Beiting-Parrish, Magdalen
%S Proceedings of the Artificial Intelligence in Measurement and Education Conference (AIME-Con): Full Papers
%D 2026
%8 October
%I National Council on Measurement in Education (NCME)
%C Wyndham Grand Pittsburgh Downtown, Pittsburgh, Pennsylvania, United States
%@ 979-8-9983004-0-0
%F jong-koo-2026-including
%X This study measures the effectiveness of including images to predict discrimination parameters of reading items. The results suggest that providing images to language models is beneficial. By using both image and text, Lasso regressions (Tibshirani, 1996) could explain data better and random forests (Breiman, 2001) showed higher accuracy.
%U https://aclanthology.org/2026.aimecon-main.71/
%P 631-638
Markdown (Informal)
[Does Including Images Improve Multimodal Language Models’ Accuracy on Item Discrimination Estimation](https://aclanthology.org/2026.aimecon-main.71/) (Jong & Koo, AIME-Con 2026)
ACL