<?xml version="1.0" encoding="UTF-8"?>
<oai_dc:dc xmlns:oai_dc="http://www.openarchives.org/OAI/2.0/oai_dc/" xmlns:dc="http://purl.org/dc/elements/1.1/" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" xsi:schemaLocation="http://www.openarchives.org/OAI/2.0/oai_dc/ http://www.openarchives.org/OAI/2.0/oai_dc.xsd">
  <dc:title>Statistical and Psychometric Evaluation of AI Systems</dc:title>
  <dc:title>R package aiEvalR version 0.1.0</dc:title>
  <dc:description>Evaluates artificial intelligence (AI) systems as measurement instruments using psychometric methods. Provides multi-facet generalizability theory (G-study and D-study) via 'lme4', reliability via the intraclass correlation coefficient (ICC), calibration via the expected calibration error (ECE) and Brier score, robustness stress testing, and group disparity diagnostics. Item-level differential item functioning (DIF) based on item response theory (IRT) is delegated to the 'aiDIF' package. Methods follow Cronbach, Gleser, Nanda and Rajaratnam (1972, &lt;ISBN:9780471188506&gt;) and Brennan (2001) &lt;doi:10.1007/978-1-4757-3456-0&gt;.</dc:description>
  <dc:type>Software</dc:type>
  <dc:relation>Depends: R (&gt;= 4.1.0)</dc:relation>
  <dc:relation>Imports: stats, utils</dc:relation>
  <dc:relation>Suggests: testthat (&gt;= 3.0.0), knitr, rmarkdown, covr, boot, lme4,
dplyr, ggplot2, aiDIF, spelling</dc:relation>
  <dc:creator>Subir Hait &lt;haitsubi@msu.edu&gt;</dc:creator>
  <dc:publisher>Comprehensive R Archive Network (CRAN)</dc:publisher>
  <dc:contributor>Subir Hait [aut, cre] (ORCID: &lt;https://orcid.org/0009-0004-9871-9677&gt;)</dc:contributor>
  <dc:rights>MIT + file LICENSE (https://CRAN.R-project.org/package=aiEvalR/LICENSE)</dc:rights>
  <dc:date>2026-08-30</dc:date>
  <dc:format>application/tgz</dc:format>
  <dc:identifier>https://CRAN.R-project.org/package=aiEvalR</dc:identifier>
  <dc:identifier>doi:10.32614/CRAN.package.aiEvalR</dc:identifier>
  <dc:language>en-US</dc:language>
</oai_dc:dc>
