<?xml version="1.0" encoding="UTF-8"?><!DOCTYPE article PUBLIC "-//NLM//DTD Journal Publishing DTD v2.0 20040830//EN" "journalpublishing.dtd"><article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" dtd-version="2.0" xml:lang="en" article-type="review-article"><front><journal-meta><journal-id journal-id-type="nlm-ta">JMIR Ment Health</journal-id><journal-id journal-id-type="publisher-id">mental</journal-id><journal-id journal-id-type="index">16</journal-id><journal-title>JMIR Mental Health</journal-title><abbrev-journal-title>JMIR Ment Health</abbrev-journal-title><issn pub-type="epub">2368-7959</issn><publisher><publisher-name>JMIR Publications</publisher-name><publisher-loc>Toronto, Canada</publisher-loc></publisher></journal-meta><article-meta><article-id pub-id-type="publisher-id">v13i1e93307</article-id><article-id pub-id-type="doi">10.2196/93307</article-id><article-categories><subj-group subj-group-type="heading"><subject>Review</subject></subj-group></article-categories><title-group><article-title>Optimizing Treatment Strategies in the Bipolar Disorder Spectrum With Classical AI Approaches: Systematic Review of Performance, Bias, and Clinical Applicability</article-title></title-group><contrib-group><contrib contrib-type="author" corresp="yes" equal-contrib="yes"><name name-style="western"><surname>De Francesco</surname><given-names>Silvia</given-names></name><degrees>MSc</degrees><xref ref-type="aff" rid="aff1">1</xref><xref ref-type="fn" rid="equal-contrib1">*</xref></contrib><contrib contrib-type="author" equal-contrib="yes"><name name-style="western"><surname>Archetti</surname><given-names>Damiano</given-names></name><degrees>MSc</degrees><xref ref-type="aff" rid="aff1">1</xref><xref ref-type="fn" rid="equal-contrib1">*</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Baronio</surname><given-names>Cesare Michele</given-names></name><degrees>PhD</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Demaria</surname><given-names>Claudio</given-names></name><degrees>MSc</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Boccali</surname><given-names>Alberto</given-names></name><degrees>MSc</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Crema</surname><given-names>Claudio</given-names></name><degrees>PhD</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Tura</surname><given-names>Giovanni Battista</given-names></name><degrees>MD</degrees><xref ref-type="aff" rid="aff2">2</xref></contrib><contrib contrib-type="author"><name name-style="western"><surname>Redolfi</surname><given-names>Alberto</given-names></name><degrees>PhD</degrees><xref ref-type="aff" rid="aff1">1</xref></contrib></contrib-group><aff id="aff1"><institution>Laboratory of Neuroinformatics, IRCCS Istituto Centro San Giovanni di Dio Fatebenefratelli</institution><addr-line>via Pilastroni 4</addr-line><addr-line>Brescia</addr-line><country>Italy</country></aff><aff id="aff2"><institution>Psychiatry Unit, IRCCS Istituto Centro San Giovanni di Dio Fatebenefratelli</institution><addr-line>Brescia</addr-line><country>Italy</country></aff><contrib-group><contrib contrib-type="editor"><name name-style="western"><surname>Morton</surname><given-names>Emma</given-names></name></contrib></contrib-group><contrib-group><contrib contrib-type="reviewer"><name name-style="western"><surname>Goodman</surname><given-names>Marianne</given-names></name></contrib><contrib contrib-type="reviewer"><name name-style="western"><surname>Vidal</surname><given-names>Nathan</given-names></name></contrib><contrib contrib-type="reviewer"><name name-style="western"><surname>Banerjee</surname><given-names>Somnath</given-names></name></contrib></contrib-group><author-notes><corresp>Correspondence to Silvia De Francesco, MSc, Laboratory of Neuroinformatics, IRCCS Istituto Centro San Giovanni di Dio Fatebenefratelli, via Pilastroni 4, Brescia, 25125, Italy; <email>sdefrancesco@fatebenefratelli.eu</email></corresp><fn fn-type="equal" id="equal-contrib1"><label>*</label><p>these authors contributed equally</p></fn></author-notes><pub-date pub-type="collection"><year>2026</year></pub-date><pub-date pub-type="epub"><day>21</day><month>7</month><year>2026</year></pub-date><volume>13</volume><elocation-id>e93307</elocation-id><history><date date-type="received"><day>11</day><month>02</month><year>2026</year></date><date date-type="rev-recd"><day>21</day><month>04</month><year>2026</year></date><date date-type="accepted"><day>21</day><month>04</month><year>2026</year></date></history><copyright-statement>&#x00A9; Silvia De Francesco, Damiano Archetti, Cesare Michele Baronio, Claudio Demaria, Alberto Boccali, Claudio Crema, Giovanni Battista Tura, Alberto Redolfi. Originally published in JMIR Mental Health (<ext-link ext-link-type="uri" xlink:href="https://mental.jmir.org">https://mental.jmir.org</ext-link>), 21.7.2026. </copyright-statement><copyright-year>2026</copyright-year><license license-type="open-access" xlink:href="https://creativecommons.org/licenses/by/4.0/"><p>This is an open-access article distributed under the terms of the Creative Commons Attribution License (<ext-link ext-link-type="uri" xlink:href="https://creativecommons.org/licenses/by/4.0/">https://creativecommons.org/licenses/by/4.0/</ext-link>), which permits unrestricted use, distribution, and reproduction in any medium, provided the original work, first published in JMIR Mental Health, is properly cited. The complete bibliographic information, a link to the original publication on <ext-link ext-link-type="uri" xlink:href="https://mental.jmir.org/">https://mental.jmir.org/</ext-link>, as well as this copyright and license information must be included.</p></license><self-uri xlink:type="simple" xlink:href="https://mental.jmir.org/2026/1/e93307"/><abstract><sec><title>Background</title><p>Bipolar disorder (BD) is a complex and heterogeneous psychiatric condition, characterized by fluctuating clinical courses that affect approximately 1%&#x2010;2% of the global population in their lifetime. Despite pharmacological advances, treatment response varies significantly among patients, making the identification of individualized treatment strategies a major challenge. Artificial Intelligence (AI), through its classical approaches, has emerged as a powerful tool in precision psychiatry to identify subtle patterns in complex data and inform personalized clinical decisions.</p></sec><sec><title>Objective</title><p>The present systematic review aimed to examine the current evidence on classical AI-supported treatment optimization in the BD spectrum.</p></sec><sec sec-type="methods"><title>Methods</title><p>The review was conducted in accordance with the PRISMA (Preferred Reporting Items for Systematic Reviews and Meta-Analyses) 2020 guidelines. Four databases (PubMed, Web of Science, Scopus, and Embase) were searched for original studies published after 2015 on the application of classical AI in the treatment of BD in adult patients. Publication bias was evaluated by visual inspection of a funnel plot. The methodological quality, risk of bias, and clinical applicability of the predictive models were assessed using the Prediction Model Risk Of Bias Assessment Tool for prediction models using regression or AI methods (PROBAST+AI; PROBAST+AI Working Group) tool.</p></sec><sec sec-type="results"><title>Results</title><p>A total of 35 studies were included and classified into 5 outcome-based categories, including acute symptomatic response, long-term maintenance response, relapse and readmission risk, safety and dose optimization, and brain aging and phenotyping. Acute symptomatic response models performed modestly (pooled area under the curve [AUC] 0.68), while imaging improved accuracy (74%&#x2010;77%). Long-term maintenance response models showed moderate-to-high performance (pooled AUC 0.80), with biomarker- and cellular-based models reaching 96%&#x2010;99% accuracy. Relapse and readmission prediction achieved a pooled AUC of 0.71, with digital phenotyping and rule-based methods performing best (AUC 0.85&#x2010;0.88). Safety and dose optimization models achieved 85%&#x2010;97% accuracy. Brain aging and phenotyping studies highlighted accelerated brain aging in BD, partially mitigated by lithium, and revealed novel data-driven subgroups. However, 3 studies were considered at high risk of bias due to small sample sizes associated with disproportionately high-performance estimates. An additional study was identified as potentially biased because it lay markedly distant from the funnel plot&#x2019;s confidence line. Finally, the PROBAST+AI assessment revealed a high risk of bias in most studies, primarily due to data analysis limitations, small sample sizes, and lack of external validation.</p></sec><sec sec-type="conclusions"><title>Conclusions</title><p>The adoption of classical AI tools in BD serves as a driver for therapeutic optimization, although current AI tools in BD should still be considered exploratory rather than ready for clinical use. Effective implementation in real-world clinical scenarios requires more robust, transparent, and externally validated models to ensure reliability and generalizability.</p></sec></abstract><kwd-group><kwd>bipolar disorder</kwd><kwd>treatment</kwd><kwd>artificial intelligence</kwd><kwd>Prediction model Risk Of Bias Assessment Tool for prediction models using regression or AI methods</kwd><kwd>PROBAST-AI</kwd><kwd>psychiatry</kwd></kwd-group></article-meta></front><body><sec id="s1" sec-type="intro"><title>Introduction</title><p>Bipolar disorder (BD) is one of the most complex and heterogeneous psychiatric conditions, characterized by recurrent mood episodes, fluctuating clinical trajectories, and significant functional impairment [<xref ref-type="bibr" rid="ref1">1</xref>]. The lifetime prevalence of BD is approximately 1%&#x2010;2% globally. By subtype, BD-I affects about 0.6% of the population and BD-II about 0.4%, with 12-month prevalence estimates of approximately 0.4%-0.3%, respectively. BD ranks among the leading causes of disability [<xref ref-type="bibr" rid="ref2">2</xref>]. Despite significant advances in pharmacological and psychosocial interventions, treatment response varies widely across patients and stages. This variability reflects the multifactorial etiology of BD, involving genetic, neurobiological, psychological, and environmental determinants [<xref ref-type="bibr" rid="ref3">3</xref>]. Identifying individualized treatment strategies remains a major challenge in clinical practice.</p><p>In addition to psychoeducation for maintaining mood stability, cognitive abilities, and social functioning, one or more drugs are usually prescribed to individuals with BD, according to the severity of each clinical situation [<xref ref-type="bibr" rid="ref4">4</xref>,<xref ref-type="bibr" rid="ref5">5</xref>]. Pharmacological management represents the first line of treatment, with commonly used mood stabilizers such as lithium and valproate [<xref ref-type="bibr" rid="ref6">6</xref>] being highly effective for both acute manic episodes and relapse prevention. Anticonvulsants such as valproate are widely used for the control of mania and maintenance, while lamotrigine is specifically indicated for the prevention of depressive relapse [<xref ref-type="bibr" rid="ref7">7</xref>]. Second-generation antipsychotics have become mainstays of therapy; quetiapine is recommended as a first-line choice for depressive episodes in BD [<xref ref-type="bibr" rid="ref8">8</xref>], olanzapine is frequently used for acute stabilization and maintenance, and lurasidone is specifically approved for acute bipolar depression [<xref ref-type="bibr" rid="ref9">9</xref>]. Therapeutic strategies must be designed to take into account several complex aspects of the illness, such as the alternation of euthymic, depressed, manic and hypomanic phases, interepisodic symptoms, and mood relapses. Despite the variety of options, the clinical management of BD faces significant challenges. The high heterogeneity of clinical phenotypes makes it difficult to predict individual response, leading to a trial-and-error approach that can delay achieving effective stabilization by up to 10 years [<xref ref-type="bibr" rid="ref10">10</xref>]. Furthermore, the risk of serious side effects, such as renal failure with lithium or cardiometabolic comorbidities with antipsychotics, requires careful balancing of therapeutic benefits against the risks of long-term toxicity [<xref ref-type="bibr" rid="ref11">11</xref>,<xref ref-type="bibr" rid="ref12">12</xref>]. To manage the chronic course of the illness and to increase the response rate of the therapy, a combination of mood stabilizers and atypical antipsychotics is usually prescribed [<xref ref-type="bibr" rid="ref13">13</xref>].</p><p>Classical AI, including machine learning (ML) and deep learning (DL), has emerged as a powerful approach within precision psychiatry over the past decade. These techniques leverage large-scale and, in some cases, multimodal data, including clinical assessments, neuroimaging, genetics, biomarkers, speech features, behavioral patterns, and digital phenotyping, to detect subtle patterns that may escape traditional statistical methods [<xref ref-type="bibr" rid="ref14">14</xref>-<xref ref-type="bibr" rid="ref16">16</xref>]. Because they are well suited for capturing nonlinear associations, handling high-dimensional and collinear data, modeling complex interactions, and generating individualized predictions, classical AI-based models have shown promise in forecasting relapse, predicting treatment response, stratifying patients into clinically meaningful subtypes, and informing individualized therapeutic decisions [<xref ref-type="bibr" rid="ref17">17</xref>,<xref ref-type="bibr" rid="ref18">18</xref>].</p><p>Methods such as ML, DL, and natural language processing (NLP; detailed definitions of these terms can be found in the <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>) have been increasingly applied to BD and related conditions. For example, ML classifiers trained on neuroimaging data have demonstrated potential in distinguishing BD from other mood disorders [<xref ref-type="bibr" rid="ref19">19</xref>], while smartphone-based digital phenotyping combined with ML has enabled the early detection of mood episodes and relapse risk [<xref ref-type="bibr" rid="ref20">20</xref>]. DL models applied to speech, actigraphy, and physiological signals have shown encouraging results for monitoring mood fluctuations and forecasting clinical deterioration [<xref ref-type="bibr" rid="ref21">21</xref>,<xref ref-type="bibr" rid="ref22">22</xref>].</p><p>Despite numerous studies using various forms of classical AI in BD research, the translation of these findings into real-world clinical practice remains limited. Challenges include the lack of standardized evaluation metrics and methodological concerns. In particular, data quality issues&#x2014;such as missing data, heterogeneity in data acquisition protocols, and variability in diagnostic criteria&#x2014;can affect model performance and generalizability. Additionally, many studies rely on limited sample sizes, increasing the risk of overfitting and reducing the robustness of the findings. Poor reproducibility remains a critical issue, as models are often not independently validated and codes or trained models are rarely shared, hindering external replication and validation across different populations [<xref ref-type="bibr" rid="ref14">14</xref>,<xref ref-type="bibr" rid="ref23">23</xref>]. Another major barrier is the perceived &#x201C;black box&#x201D; nature of many ML approaches that leads to poor interpretability of models and complicates clinical trust and regulatory approval. Furthermore, ethical, legal, and privacy considerations surrounding AI-driven mental health interventions require critical analysis before widespread implementation [<xref ref-type="bibr" rid="ref24">24</xref>-<xref ref-type="bibr" rid="ref26">26</xref>].</p><p>In this context of growing interest in AI applications in BD, this systematic review provides a structured synthesis of current evidence on data-driven treatment optimization. This review maps the range of AI methodologies applied to pharmacological interventions across the BD spectrum, followed by an evaluation of their clinical objectives, including treatment-response prediction and decision support for therapeutic planning. This systematic review aims to identify and compare evidence on the use of classical AI for treatment optimization in BD and to assess methodological quality, validation strategies, and risk of bias using the Prediction model Risk Of Bias Assessment Tool for prediction models using regression or AI methods (PROBAST+AI) framework [<xref ref-type="bibr" rid="ref27">27</xref>]. We highlighted strengths, limitations, and persistent gaps in the current literature to inform future research toward clinically robust and scalable models for real-world implementation.</p></sec><sec id="s2" sec-type="methods"><title>Methods</title><sec id="s2-1"><title>Protocol</title><p>This systematic review was conducted in accordance with the PRISMA (Preferred Reporting Items for Systematic Reviews and Meta-Analyses; <xref ref-type="supplementary-material" rid="app2">Checklist 1</xref>) 2020 guidelines.</p></sec><sec id="s2-2"><title>Eligibility Criteria</title><p>Studies were eligible if they: (1) reported original research applying classical AI, ML, DL, or hybrid models; (2) focused on treatment of individuals with BD; (3) focused on adult populations (aged &#x2265;18 years) or included mixed samples with both adults and adolescents; (4) were written in English; and (5) were published after 2015.</p><p>Studies were excluded if they met at least one of the following criteria: (1) were reviews, meta-analyses, or commentaries; (2) were clinical trial protocols; (3) were conference abstracts, editorials, letters, lectures, or book chapters; (4) did not include original data; and (5) exclusively enrolled populations aged &#x003C;18 years.</p></sec><sec id="s2-3"><title>Search Strategy and Output Assessment</title><p>A systematic literature search was conducted across 4 electronic databases, including PubMed, Web of Science, Scopus, and Embase. The searches were performed on November 17, 2025. Queries are detailed in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>. All retrieved records were imported into the &#x201C;Rayyan&#x201D; (Rayyan Systems, Inc) application [<xref ref-type="bibr" rid="ref28">28</xref>] and duplicates were removed.</p><p>Following deduplication, records were screened in 2 sequential stages. In the first stage, titles and abstracts were independently assessed by 3 reviewers (SDF, DA, and AR) to identify studies potentially eligible based on the predefined inclusion and exclusion criteria. In the second stage, full texts of the selected articles were retrieved and independently evaluated for eligibility by the same reviewers. Any uncertainties arising during either stage were resolved through discussion, and when necessary, a consensus meeting was convened to reach a final decision regarding study inclusion or exclusion.</p><p>In addition, the annual trend of publications was examined to assess changes in research activity over time.</p></sec><sec id="s2-4"><title>Data Extraction</title><p>From each study, the following information was extracted:</p><list list-type="bullet"><list-item><p>Aim of the study (respect to treatment, prediction, or clinical decision support)</p></list-item><list-item><p>Sample size of BD group</p></list-item><list-item><p>Drugs</p></list-item><list-item><p>Classical AI method used</p></list-item><list-item><p>Biomarker domain</p></list-item><list-item><p>Validation strategy</p></list-item><list-item><p>Outcomes and performance metrics</p></list-item><list-item><p>Relevance</p></list-item></list><p>Studies were grouped into five categories based on their primary clinical aim: (1) acute symptomatic response, referring to the prediction of short-term (4/12 weeks) treatment success based on symptom severity reduction; (2) long-term (more than 6 months) maintenance response, referring to the ability to remain episode-free over time; (3) relapse and readmission risk, including the prediction of manic and depressive relapse as well as hospital readmission; (4) safety and dose optimization, focused on treatment personalization through dose adjustment and adverse-event prevention; and (5) brain aging and phenotyping aimed at identifying biologically or clinically meaningful patient subgroups.</p><p>This framework was designed to capture the main clinical areas in BD in which AI provides practical benefit.</p><p>Studies were further categorized according to the type of biomarkers included: only clinical data (clinical); genomic and clinical data (genomic); imaging and clinical data (imaging); wearable devices (including smartphone sensors); and clinical data (wearable). They were also sorted according to the modeling approach used (ie, logistic regression, Na&#x00EF;ve Bayes, decision trees, random forest, support vector machines (SVM), gradient boosting, and neural networks).</p><p>All studies were classified into 3 tiers of validation strategy. Tier 1 included studies where validation on an independent data cohort was performed, Tier 2 were those studies where only internal validation strategies (eg, cross-validation or leave-one-out approaches) were performed, and Tier 3 were exploratory studies.</p><p>The performance metrics considered included accuracy, area under the curve (AUC), sensitivity, specificity, and their CIs computed on the validation or test sets (when available or derivable).</p></sec><sec id="s2-5"><title>Bias Assessment Risk</title><p>To explore potential publication bias, we constructed a funnel plot. The plot displays the best AUC for each study on the X-axis and the inverse square root of the sample size on the Y-axis. This approach allowed visualization of the relationship between model performance and study size to identify potentially biased studies.</p><p>The assessment of methodological quality, risk of bias, and applicability of the included predictive models was performed using the PROBAST+AI method [<xref ref-type="bibr" rid="ref27">27</xref>]. The evaluation process was structured into 2 main components, model development, which analyzed the quality and generalizability of the algorithm&#x2019;s construction, and model evaluation, which assessed the credibility and clinical applicability of performance estimates that may be affected by systematic bias. The analysis was structured into 4 fundamental domains, including participants and data sources, predictors, outcomes, and analysis. An additional overall judgment was provided following the PROBAST+AI guidelines.</p></sec><sec id="s2-6"><title>Statistical Analysis</title><p>To assess the average performance of the models stratified by the defined categories (ie, acute symptomatic response, long-term maintenance response, relapse and readmission risk, safety and dose optimization, and brain aging and phenotyping), a pooled AUC was calculated from study-specific AUCs and their 95% CIs (when available) using a random-effects meta-analysis based on the DerSimonian-Laird method [<xref ref-type="bibr" rid="ref29">29</xref>]. The weights were calculated based on the 95% CIs for each study&#x2019;s AUC. Heterogeneity was summarized using the inconsistency index (<italic>I</italic><sup>2</sup>) and the between-study variance (&#x03C4;<sup>2</sup>) measure.</p><p>To further evaluate potential publication bias, the Egger test was performed on logit-transformed AUC values. This statistical test assesses asymmetry in the funnel plot, indicating whether smaller studies tend to report systematically higher or lower performance estimates compared to larger studies. The logit transformation was applied to stabilize variance across studies.</p><p>All analyses were performed using Python (version 3.8.8; Python Software Foundation).</p></sec></sec><sec id="s3" sec-type="results"><title>Results</title><sec id="s3-1"><title>Study Selection</title><p>The PubMed query returned 127 articles; the SCOPUS query returned 297 articles; the Web of Science query returned 172 articles; the Embase query returned 119 articles. After removing duplicates, a sample of 327 studies was screened; 35 studies met all the eligibility criteria and were included in the final synthesis (<xref ref-type="fig" rid="figure1">Figure 1</xref>).</p><fig position="float" id="figure1"><label>Figure 1.</label><caption><p>PRISMA (Preferred Reporting Items for Systematic Reviews and Meta-Analyses) 2020 flow diagram illustrating the paper selection process.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="mental_v13i1e93307_fig01.png"/></fig></sec><sec id="s3-2"><title>Trends in Scientific Publications</title><p>Over the past decade, scientific interest in applying AI to the treatment and long-term management of BD has grown substantially. The publication trend of the 327 records screened in this review, normalized on the total number of publications per year in PubMed, follows an exponential trajectory (<italic>R</italic><sup>2</sup>=0.979); while contributions in the early 2010s were relatively limited, the volume of research has steadily increased, with particularly pronounced growth in the past 6 years (<xref ref-type="fig" rid="figure2">Figure 2</xref>).</p><fig position="float" id="figure2"><label>Figure 2.</label><caption><p>Annual number of publications on AI applications in the treatment of bipolar disorder, normalized by the total number of PubMed publications per year and scaled by 10&#x2076;.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="mental_v13i1e93307_fig02.png"/></fig></sec><sec id="s3-3"><title>Characteristics of Included Studies</title><p>The final set comprised 35 [<xref ref-type="bibr" rid="ref22">22</xref>,<xref ref-type="bibr" rid="ref30">30</xref>-<xref ref-type="bibr" rid="ref63">63</xref>] studies published between 2015 and 2025. Most studies were conducted in the United States, China, the European Union (EU), Canada, and the United Kingdom, with sample sizes ranging from about 10 to approximately 530 thousand participants. A detailed summary of study characteristics is provided in <xref ref-type="table" rid="table1">Table 1</xref>. The aims of the studies are reported in Table S1 in the <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>. Data sources varied widely and included different clinical assessment tools and rating scales (eg, Young Mania Rating Scale [YMRS] and Hamilton Depression Scale [HAMD]), electronic health records (EHR), neuroimaging data, digital phenotyping (smartphone sensors and wearables), and multimodal datasets. AI methodologies included classical ML models (eg, random forests [RF], SVMs, DL architectures (convolutional neural networks [CNNs] and recurrent neural networks [RNNs]), NLP approaches, and hybrid computational frameworks. Based on the validation strategy, we classified most studies [<xref ref-type="bibr" rid="ref32">32</xref>-<xref ref-type="bibr" rid="ref34">34</xref>,<xref ref-type="bibr" rid="ref36">36</xref>,<xref ref-type="bibr" rid="ref37">37</xref>,<xref ref-type="bibr" rid="ref46">46</xref>-<xref ref-type="bibr" rid="ref49">49</xref>,<xref ref-type="bibr" rid="ref51">51</xref>-<xref ref-type="bibr" rid="ref54">54</xref>,<xref ref-type="bibr" rid="ref56">56</xref>,<xref ref-type="bibr" rid="ref58">58</xref>-<xref ref-type="bibr" rid="ref63">63</xref>] (n=20, 57%) as Tier 2, a significant proportion [<xref ref-type="bibr" rid="ref22">22</xref>,<xref ref-type="bibr" rid="ref30">30</xref>,<xref ref-type="bibr" rid="ref31">31</xref>,<xref ref-type="bibr" rid="ref35">35</xref>,<xref ref-type="bibr" rid="ref38">38</xref>,<xref ref-type="bibr" rid="ref41">41</xref>-<xref ref-type="bibr" rid="ref45">45</xref>,<xref ref-type="bibr" rid="ref50">50</xref>,<xref ref-type="bibr" rid="ref55">55</xref>,<xref ref-type="bibr" rid="ref57">57</xref>] (n=13, 37%) as Tier 3, while only 2 studies [<xref ref-type="bibr" rid="ref39">39</xref>,<xref ref-type="bibr" rid="ref40">40</xref>] were categorized as Tier 1.</p><table-wrap id="t1" position="float"><label>Table 1.</label><caption><p>Characteristics of the 35 included studies (2015&#x2010;2025) on AI applications in optimizing bipolar disorder treatment.</p></caption><table id="table1" frame="hsides" rules="groups"><thead><tr><td align="left" valign="bottom">Study</td><td align="left" valign="bottom">Sample size</td><td align="left" valign="bottom">Category</td><td align="left" valign="bottom">AI models</td><td align="left" valign="bottom">Validation tier</td><td align="left" valign="bottom">Code availability</td><td align="left" valign="bottom">Drugs</td><td align="left" valign="bottom">Outcomes</td></tr></thead><tbody><tr><td align="left" valign="top">Agniel et al (2024) [<xref ref-type="bibr" rid="ref30">30</xref>]</td><td align="left" valign="top">BD<sup><xref ref-type="table-fn" rid="table1fn1">a</xref></sup> (n=5500) and others</td><td align="left" valign="top">4</td><td align="left" valign="top">Linear models; tree-based models; neural networks</td><td align="left" valign="top">Tier 3</td><td align="left" valign="top">No</td><td align="left" valign="top">Atypical antipsychotics</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>RMST:<sup><xref ref-type="table-fn" rid="table1fn2">b</xref></sup> 21.9-22.1</p></list-item><list-item><p>Diabetes risk at 24 months: 6.3%-7.1%</p></list-item><list-item><p>RR:<sup><xref ref-type="table-fn" rid="table1fn3">c</xref></sup> 0.9</p></list-item></list></td></tr><tr><td align="left" valign="top">Brodeur et al (2021) [<xref ref-type="bibr" rid="ref31">31</xref>]</td><td align="left" valign="top">BD I (n=865), BD II (n=764), BD NOS<sup><xref ref-type="table-fn" rid="table1fn4">d</xref></sup> (n=166)</td><td align="left" valign="top">5</td><td align="left" valign="top">Clustering models</td><td align="left" valign="top">Tier 3</td><td align="left" valign="top">No</td><td align="left" valign="top">Mood stabilizers, mood-stabilizing anticonvulsants, atypical antipsychotics, typical antipsychotics, antidepressant, anxiolytic</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>QIDS-SR16<sup><xref ref-type="table-fn" rid="table1fn5">e</xref></sup>=6-7</p></list-item><list-item><p>MADRS<sup><xref ref-type="table-fn" rid="table1fn6">f</xref></sup>=4-6</p></list-item><list-item><p>YMRS<sup><xref ref-type="table-fn" rid="table1fn7">g</xref></sup>=0</p></list-item><list-item><p>FAST<sup><xref ref-type="table-fn" rid="table1fn8">h</xref></sup>=10-15</p></list-item><list-item><p>CGI-S<sup><xref ref-type="table-fn" rid="table1fn9">i</xref></sup>=2-3</p></list-item><list-item><p>GAF<sup><xref ref-type="table-fn" rid="table1fn10">j</xref></sup>=45.2%-59.2%</p></list-item></list></td></tr><tr><td align="left" valign="top">Cearns et al (2022) [<xref ref-type="bibr" rid="ref32">32</xref>]</td><td align="left" valign="top">BD I (n=803), BD II (n=203), BD Schizoaffective (n=18), BD III (n=3), BD NOS (n=7)</td><td align="left" valign="top">2</td><td align="left" valign="top">Linear models; tree-based models</td><td align="left" valign="top">Tier 2</td><td align="left" valign="top">No</td><td align="left" valign="top">Mood stabilizers</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>R&#x00B2;=1.8-13.7%</p></list-item><list-item><p>Balanced accuracy=58.9-63.7%</p></list-item></list></td></tr><tr><td align="left" valign="top">Diaz-Zuluaga et al (2023) [<xref ref-type="bibr" rid="ref33">33</xref>]</td><td align="left" valign="top">BD I, BD II, (n total=172)</td><td align="left" valign="top">2</td><td align="left" valign="top">Tree-based models</td><td align="left" valign="top">Tier 2</td><td align="left" valign="top">No</td><td align="left" valign="top">Mood stabilizers</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>Sensitivity=50%-84.6%</p></list-item><list-item><p>Specificity=93.8%-94.5%</p></list-item><list-item><p>AUC=72.2%-89.2%</p></list-item><list-item><p>Accuracy=87.8%-92.4%</p></list-item></list></td></tr><tr><td align="left" valign="top">Edgcomb et al (2021) [<xref ref-type="bibr" rid="ref34">34</xref>]</td><td align="left" valign="top">BD (n=502) and others</td><td align="left" valign="top">3</td><td align="left" valign="top">Tree-based models</td><td align="left" valign="top">Tier 2</td><td align="left" valign="top">No</td><td align="left" valign="top">Mood stabilizers, antidepressants, anxiolytics, antipsychotics</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>AUC=86%</p></list-item><list-item><p>Sensitivity=82%</p></list-item><list-item><p>Specificity=80%</p></list-item><list-item><p>Accuracy=80%</p></list-item></list></td></tr><tr><td align="left" valign="top">Fleck et al (2017) [<xref ref-type="bibr" rid="ref35">35</xref>]</td><td align="left" valign="top">BD I (n=20)</td><td align="left" valign="top">1</td><td align="left" valign="top">Tree-based models; kernel-based models</td><td align="left" valign="top">Tier 3</td><td align="left" valign="top">No</td><td align="left" valign="top">Mood stabilizers</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>Accuracy=50%-100%</p></list-item><list-item><p>YMRS reduction accuracy=60.99%-91.77%</p></list-item></list></td></tr><tr><td align="left" valign="top">Hayes et al (2024) [<xref ref-type="bibr" rid="ref36">36</xref>]</td><td align="left" valign="top">BD (n=31,518)</td><td align="left" valign="top">2</td><td align="left" valign="top">Linear models; tree-based models; Bayesian models</td><td align="left" valign="top">Tier 2</td><td align="left" valign="top">Available [<xref ref-type="bibr" rid="ref64">64</xref>]</td><td align="left" valign="top">Mood stabilizers, atypical antipsychotics</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>Accuracy=50%-61.6%</p></list-item></list></td></tr><tr><td align="left" valign="top">Kim et al (2019) [<xref ref-type="bibr" rid="ref37">37</xref>]</td><td align="left" valign="top">BD I, BD II, (n total=482)</td><td align="left" valign="top">1</td><td align="left" valign="top">Linear models</td><td align="left" valign="top">Tier 2</td><td align="left" valign="top">No</td><td align="left" valign="top">Mood stabilizers, atypical antipsychotics</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>R&#x00B2; =17.4%-32.1%</p></list-item></list></td></tr><tr><td align="left" valign="top">Lei et al (2022) [<xref ref-type="bibr" rid="ref38">38</xref>]</td><td align="left" valign="top">BD (n=109) and others</td><td align="left" valign="top">1</td><td align="left" valign="top">Kernel-based models</td><td align="left" valign="top">Tier 3</td><td align="left" valign="top">No</td><td align="left" valign="top">Mood stabilizers, atypical antipsychotics</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>Accuracy=58%-74%</p></list-item><list-item><p>Sensitivity=57%-75%</p></list-item><list-item><p>Specificity=42%-91%</p></list-item></list></td></tr><tr><td align="left" valign="top">Lieslehto et al (2025) [<xref ref-type="bibr" rid="ref39">39</xref>]</td><td align="left" valign="top">FEBD<sup><xref ref-type="table-fn" rid="table1fn11">k</xref></sup> (n=44,969)</td><td align="left" valign="top">3</td><td align="left" valign="top">Tree-based models</td><td align="left" valign="top">Tier 1</td><td align="left" valign="top">Available [<xref ref-type="bibr" rid="ref65">65</xref>]</td><td align="left" valign="top">Mood stabilizers, antipsychotics, antidepressants, anxiolytic</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>HR<sup><xref ref-type="table-fn" rid="table1fn12">l</xref></sup>=0.42-1.50</p></list-item><list-item><p>AUC<sup><xref ref-type="table-fn" rid="table1fn13">m</xref></sup>=71%-77%</p></list-item><list-item><p>Sensitivity=75.8%-75.9%</p></list-item><list-item><p>Specificity=59.7%-65.8%</p></list-item></list></td></tr><tr><td align="left" valign="top">Lieslehto et al (2025) [<xref ref-type="bibr" rid="ref40">40</xref>]</td><td align="left" valign="top">FEBD (n=44,192)</td><td align="left" valign="top">3</td><td align="left" valign="top">Linear models; tree-based models; kernel-based models</td><td align="left" valign="top">Tier 1</td><td align="left" valign="top">Available [<xref ref-type="bibr" rid="ref66">66</xref>]</td><td align="left" valign="top">Mood stabilizers, antipsychotics, antidepressants, anxiolytic</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>AUC = 65%-85%</p></list-item><list-item><p>Brier score=0.08-0.12</p></list-item><list-item><p>Calibration slope=0.82-0.95</p></list-item><list-item><p>Sensitivity=56.08%-61.82%</p></list-item><list-item><p>Specificity=70.83%-72.80%</p></list-item><list-item><p>HR=3.84-4.61</p></list-item></list></td></tr><tr><td align="left" valign="top">Lv et al (2025) [<xref ref-type="bibr" rid="ref41">41</xref>]</td><td align="left" valign="top">BD (n=68) and others</td><td align="left" valign="top">1</td><td align="left" valign="top">Kernel-based models</td><td align="left" valign="top">Tier 3</td><td align="left" valign="top">No</td><td align="left" valign="top">Mood stabilizers, atypical antipsychotics, mood-stabilizing anticonvulsants, anxiolytic</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>Accuracy: 68%</p></list-item><list-item><p>Sensitivity: 79%-89%</p></list-item><list-item><p>Specificity: 59%-60%</p></list-item><list-item><p>AUC: 75%-76%</p></list-item><list-item><p>Correlation: 0.31-0.34</p></list-item></list></td></tr><tr><td align="left" valign="top">Marie-Claire et al 2022 [<xref ref-type="bibr" rid="ref42">42</xref>]</td><td align="left" valign="top">BD I (n=70)</td><td align="left" valign="top">2</td><td align="left" valign="top">Tree-based models; Linear models</td><td align="left" valign="top">Tier 3</td><td align="left" valign="top">No</td><td align="left" valign="top">Mood stabilizers</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>AUC=73%</p></list-item><list-item><p>Accuracy=84.1%</p></list-item><list-item><p>Sensitivity=97.9%</p></list-item><list-item><p>Specificity=13.3%-40%</p></list-item></list></td></tr><tr><td align="left" valign="top">Mizrahi et al (2023) [<xref ref-type="bibr" rid="ref43">43</xref>]</td><td align="left" valign="top">BD (n=43) and others</td><td align="left" valign="top">2</td><td align="left" valign="top">Linear models; tree-based models; kernel-based models; neural networks</td><td align="left" valign="top">Tier 3</td><td align="left" valign="top">Available [<xref ref-type="bibr" rid="ref67">67</xref>]</td><td align="left" valign="top">Mood stabilizers</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>AUC: 95%-99%</p></list-item><list-item><p>Accuracy: 86%-96.5%</p></list-item></list></td></tr><tr><td align="left" valign="top">Mora et al (2025) [<xref ref-type="bibr" rid="ref44">44</xref>]</td><td align="left" valign="top">BD I (n=22), BD II (n=4), BD unspecified (n=27), and others</td><td align="left" valign="top">4</td><td align="left" valign="top">Other (NLP)<sup><xref ref-type="table-fn" rid="table1fn14">n</xref></sup></td><td align="left" valign="top">Tier 3</td><td align="left" valign="top">No</td><td align="left" valign="top">Atypical antipsychotics, typical antipsychotics</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>No performances, only clinical outcomes reported as percentage change (anxiety, depressive, positive, and negative symptoms)</p></list-item></list></td></tr><tr><td align="left" valign="top">Nestsiarovich et al (2021) [<xref ref-type="bibr" rid="ref45">45</xref>]</td><td align="left" valign="top">BD (n=529,359)</td><td align="left" valign="top">3</td><td align="left" valign="top">Linear models; tree-based models</td><td align="left" valign="top">Tier 3</td><td align="left" valign="top">Available [<xref ref-type="bibr" rid="ref68">68</xref>]</td><td align="left" valign="top">Mood stabilizers, mood-stabilizing anticonvulsants, atypical antipsychotics, typical antipsychotics, antidepressants</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>HR: 0.28-2.33</p></list-item></list></td></tr><tr><td align="left" valign="top">Nielsen et al (2019) [<xref ref-type="bibr" rid="ref46">46</xref>]</td><td align="left" valign="top">BD and others (n total=84)</td><td align="left" valign="top">1</td><td align="left" valign="top">Linear models; Kernel-based models</td><td align="left" valign="top">Tier 2</td><td align="left" valign="top">No</td><td align="left" valign="top">Mood stabilizers, anticonvulsants, antidepressants, antipsychotics, anxiolytics</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>Accuracy=28.59%-58.98%</p></list-item><list-item><p>AUC=24.11%-59.89%</p></list-item></list></td></tr><tr><td align="left" valign="top">Nunes et al (2020) [<xref ref-type="bibr" rid="ref47">47</xref>]</td><td align="left" valign="top">BD I, BD II, (n total=1266)</td><td align="left" valign="top">2</td><td align="left" valign="top">Tree-based Models</td><td align="left" valign="top">Tier 2</td><td align="left" valign="top">No</td><td align="left" valign="top">Mood stabilizers</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>AUC = 60.92%-71.49%</p></list-item><list-item><p>Accuracy = 58.67%-65.23%</p></list-item><list-item><p>Sensitivity = 51.99%-66.02%</p></list-item><list-item><p>Specificity = 64.85%-64.90%</p></list-item></list></td></tr><tr><td align="left" valign="top">Palau et al (2023) [<xref ref-type="bibr" rid="ref48">48</xref>]</td><td align="left" valign="top">BD I and others (n total=78)</td><td align="left" valign="top">3</td><td align="left" valign="top">Linear Models</td><td align="left" valign="top">Tier 2</td><td align="left" valign="top">Available [<xref ref-type="bibr" rid="ref69">69</xref>]</td><td align="left" valign="top">Mood stabilizers, typical antipsychotics</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>HR: 2.35</p></list-item><list-item><p>AUC: 65%</p></list-item><list-item><p>ECE<sup><xref ref-type="table-fn" rid="table1fn15">o</xref></sup>: &#x003C;0.20</p></list-item></list></td></tr><tr><td align="left" valign="top">Pettorruso et al (2023) [<xref ref-type="bibr" rid="ref49">49</xref>]</td><td align="left" valign="top">BD (n=39) and others</td><td align="left" valign="top">1</td><td align="left" valign="top">Tree-based Models</td><td align="left" valign="top">Tier 2</td><td align="left" valign="top">No</td><td align="left" valign="top">Antidepressants</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>Accuracy=66.26%-68.6%</p></list-item></list></td></tr><tr><td align="left" valign="top">Poulos et al (2024) [<xref ref-type="bibr" rid="ref50">50</xref>]</td><td align="left" valign="top">BD I and others (n total=38,762)</td><td align="left" valign="top">4</td><td align="left" valign="top">Linear models; tree-based models; ensemble models</td><td align="left" valign="top">Tier 3</td><td align="left" valign="top">Available [<xref ref-type="bibr" rid="ref70">70</xref>]</td><td align="left" valign="top">Atypical antipsychotics, typical antipsychotics</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>ATE<sup><xref ref-type="table-fn" rid="table1fn16">p</xref></sup>=&#x2013;1.9%-2.2%</p></list-item></list></td></tr><tr><td align="left" valign="top">Ross et al (2024) [<xref ref-type="bibr" rid="ref51">51</xref>]</td><td align="left" valign="top">BD (n =34,071) and others</td><td align="left" valign="top">3</td><td align="left" valign="top">Tree-based models</td><td align="left" valign="top">Tier 2</td><td align="left" valign="top">No</td><td align="left" valign="top">ND</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>ATE=&#x2013;9.6%</p></list-item><list-item><p>ITR<sup><xref ref-type="table-fn" rid="table1fn17">q</xref></sup>=18.1%</p></list-item></list></td></tr><tr><td align="left" valign="top">Ryan et al (2022) [<xref ref-type="bibr" rid="ref52">52</xref>]</td><td align="left" valign="top">BD (n=47) and others</td><td align="left" valign="top">5</td><td align="left" valign="top">Linear models; tree-based models</td><td align="left" valign="top">Tier 2</td><td align="left" valign="top">No</td><td align="left" valign="top">Antidepressants, Mood stabilizers, Antipsychotics</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>Cohen <italic>d</italic>=0.42-0.55</p></list-item></list></td></tr><tr><td align="left" valign="top">Salari et al (2025) [<xref ref-type="bibr" rid="ref53">53</xref>]</td><td align="left" valign="top">BD (n=60)</td><td align="left" valign="top">2</td><td align="left" valign="top">Bayesian models; Tree-based models; Kernel-based models</td><td align="left" valign="top">Tier 2</td><td align="left" valign="top">No</td><td align="left" valign="top">Mood stabilizers</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>AUC = 15-86%</p></list-item><list-item><p>Sensitivity = 20-99%</p></list-item><list-item><p>Specificity = 74-85%</p></list-item><list-item><p>Precision = 35-95%</p></list-item></list></td></tr><tr><td align="left" valign="top">Salvini et al (2015) [<xref ref-type="bibr" rid="ref54">54</xref>]</td><td align="left" valign="top">BD (n=108)</td><td align="left" valign="top">3</td><td align="left" valign="top">Other (ILP)</td><td align="left" valign="top">Tier 2</td><td align="left" valign="top">No</td><td align="left" valign="top">Mood stabilizers, anticonvulsant, atypical antipsychotics, typical antipsychotics, antidepressants</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>Accuracy=85%-91%</p></list-item><list-item><p>Sensitivity=73%-92%</p></list-item><list-item><p>Precision=80%-90%</p></list-item><list-item><p>Specificity=59%-95%</p></list-item></list></td></tr><tr><td align="left" valign="top">Scott et al (2021) [<xref ref-type="bibr" rid="ref55">55</xref>]</td><td align="left" valign="top">BD (n=164)</td><td align="left" valign="top">2</td><td align="left" valign="top">Tree-based models</td><td align="left" valign="top">Tier 3</td><td align="left" valign="top">No</td><td align="left" valign="top">Mood stabilizers</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>Accuracy=90%</p></list-item><list-item><p>PPV<sup><xref ref-type="table-fn" rid="table1fn18">r</xref></sup>&#x003E;80%</p></list-item><list-item><p>NPV<sup><xref ref-type="table-fn" rid="table1fn19">s</xref></sup>&#x003E;95%</p></list-item></list></td></tr><tr><td align="left" valign="top">Stone et al (2021) [<xref ref-type="bibr" rid="ref56">56</xref>]</td><td align="left" valign="top">BD I (n=1,720), BD II (n=481), BD Schizoaffective (n=4), BD NOS (n=5)</td><td align="left" valign="top">2</td><td align="left" valign="top">Linear models; tree-based models</td><td align="left" valign="top">Tier 2</td><td align="left" valign="top">No</td><td align="left" valign="top">Mood stabilizers</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>AUC=32%-62%</p></list-item><list-item><p>Accuracy=51%-89%</p></list-item><list-item><p>F1=0%-34%</p></list-item><list-item><p>NPV=52%-89%</p></list-item><list-item><p>PPV=0%-70%</p></list-item><list-item><p>K=&#x2013;0.03-0.20</p></list-item><list-item><p>Sensitivity=0%-24%</p></list-item><list-item><p>Specificity=76%-100%</p></list-item></list></td></tr><tr><td align="left" valign="top">Tripathi et al (2023) [<xref ref-type="bibr" rid="ref57">57</xref>]</td><td align="left" valign="top">BD (n=6) and others</td><td align="left" valign="top">2</td><td align="left" valign="top">Kernel-based models; tree-based models</td><td align="left" valign="top">Tier 3</td><td align="left" valign="top">No</td><td align="left" valign="top">Mood stabilizers, mood-stabilizing anticonvulsants</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>Accuracy=60%-99%</p></list-item><list-item><p>AUC=63%-99%</p></list-item></list></td></tr><tr><td align="left" valign="top">Van Gestel et al (2019) [<xref ref-type="bibr" rid="ref58">58</xref>]</td><td align="left" valign="top">BD (n=84) and others</td><td align="left" valign="top">5</td><td align="left" valign="top">Bayesian models</td><td align="left" valign="top">Tier 2</td><td align="left" valign="top">No</td><td align="left" valign="top">Mood stabilizers</td><td align="left" valign="top">BrainAGE<sup><xref ref-type="table-fn" rid="table1fn20">t</xref></sup> =&#x2013;1.97-4.96</td></tr><tr><td align="left" valign="top">Wu et al (2025) [<xref ref-type="bibr" rid="ref22">22</xref>]</td><td align="left" valign="top">BD (n=24)</td><td align="left" valign="top">3</td><td align="left" valign="top">Linear models; Tree-based models; clustering models</td><td align="left" valign="top">Tier 3</td><td align="left" valign="top">No</td><td align="left" valign="top">Mood stabilizers, mood-stabilizing anticonvulsants, atypical antipsychotic, typical antipsychotic, antidepressant</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>Accuracy=63%-94%</p></list-item><list-item><p>AUC=51%-85%</p></list-item><list-item><p>Sensitivity=8%-88%</p></list-item><list-item><p>Specificity=57%-97%</p></list-item><list-item><p>Precision=3%-48%</p></list-item><list-item><p>F1=5%-57%</p></list-item></list></td></tr><tr><td align="left" valign="top">Xi et al (2025) [<xref ref-type="bibr" rid="ref59">59</xref>]</td><td align="left" valign="top">BD (n=580) and others</td><td align="left" valign="top">1</td><td align="left" valign="top">Kernel-based models</td><td align="left" valign="top">Tier 2</td><td align="left" valign="top">No</td><td align="left" valign="top">Atypical antipsychotic</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>HAMD<sup><xref ref-type="table-fn" rid="table1fn21">u</xref></sup> reduction ratio: r = 0.39-0.45</p></list-item><list-item><p>MSE<sup><xref ref-type="table-fn" rid="table1fn22">v</xref></sup>=0.24-0.54</p></list-item></list></td></tr><tr><td align="left" valign="top">Zhang et al (2022) [<xref ref-type="bibr" rid="ref60">60</xref>]</td><td align="left" valign="top">BD (n=236)</td><td align="left" valign="top">4</td><td align="left" valign="top">Neural networks</td><td align="left" valign="top">Tier 2</td><td align="left" valign="top">No</td><td align="left" valign="top">Mood stabilizers, anxiolytics, atypical antipsychotic, mood-stabilizing anticonvulsant</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>AC<sup><xref ref-type="table-fn" rid="table1fn23">w</xref></sup>: 0.21-1.20</p></list-item><list-item><p>PC:<sup><xref ref-type="table-fn" rid="table1fn24">x</xref></sup> 0.19-1.23</p></list-item><list-item><p>D%:<sup><xref ref-type="table-fn" rid="table1fn25">y</xref></sup> 0.00-18.52</p></list-item></list></td></tr><tr><td align="left" valign="top">Zhang et al (2025) [<xref ref-type="bibr" rid="ref61">61</xref>]</td><td align="left" valign="top">BD (n=77) and others</td><td align="left" valign="top">1</td><td align="left" valign="top">Kernel-based models</td><td align="left" valign="top">Tier 2</td><td align="left" valign="top">No</td><td align="left" valign="top">Mood stabilizers, atypical antipsychotic, anxiolytics, mood-stabilizing anticonvulsant</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>Accuracy=74.4%-76.9%</p></list-item><list-item><p>Sensitivity=72.7%-76.6%</p></list-item><list-item><p>Specificity=74.7%-78.3%</p></list-item><list-item><p>AUC=79.2%-80.3%</p></list-item><list-item><p>r<sup><xref ref-type="table-fn" rid="table1fn26">z</xref></sup>=0.662-0.722</p></list-item><list-item><p>MSE=0.021-0.022</p></list-item><list-item><p>Cohen <italic>d</italic>=0.122-0.135</p></list-item></list></td></tr><tr><td align="left" valign="top">Zheng et al (2022) [<xref ref-type="bibr" rid="ref62">62</xref>]</td><td align="left" valign="top">BD (n=177)</td><td align="left" valign="top">4</td><td align="left" valign="top">Linear models; tree-based models; kernel-based models; neural networks</td><td align="left" valign="top">Tier 2</td><td align="left" valign="top">No</td><td align="left" valign="top">Mood stabilizers, mood-stabilizing anticonvulsant, antipsychotics, antidepressants, anxiolytics</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>Accuracy=73%-85%</p></list-item><list-item><p>AUC=77%-91%</p></list-item><list-item><p>Sensitivity=78%-85%</p></list-item><list-item><p>Specificity=33%-83%</p></list-item><list-item><p>Precision=29%-96%</p></list-item><list-item><p>Recall=33%-85%</p></list-item><list-item><p>F1=31%-90%</p></list-item></list></td></tr><tr><td align="left" valign="top">Zhu et al (2022) [<xref ref-type="bibr" rid="ref63">63</xref>]</td><td align="left" valign="top">BD and others (n total=393)</td><td align="left" valign="top">4</td><td align="left" valign="top">Tree-based models</td><td align="left" valign="top">Tier 2</td><td align="left" valign="top">No</td><td align="left" valign="top">Atypical antipsychotics</td><td align="left" valign="top"><list list-type="bullet"><list-item><p>MAE:<sup><xref ref-type="table-fn" rid="table1fn27">aa</xref></sup> 0.046</p></list-item><list-item><p>MSE: 0.0036-0.0043</p></list-item><list-item><p>RMSE: 0.060-0.065</p></list-item><list-item><p>MRE:<sup><xref ref-type="table-fn" rid="table1fn28">ab</xref></sup> 11-18%</p></list-item></list></td></tr></tbody></table><table-wrap-foot><fn id="table1fn1"><p><sup>a</sup>BD: bipolar disorder.</p></fn><fn id="table1fn2"><p><sup>b</sup>RMST: restricted mean survival time.</p></fn><fn id="table1fn3"><p><sup>c</sup>RR: relative risk.</p></fn><fn id="table1fn4"><p><sup>d</sup>NOS: not otherwise specified.</p></fn><fn id="table1fn5"><p><sup>e</sup>QIDS-SR16: Quick Inventory of Depressive Symptoms 16 items.</p></fn><fn id="table1fn6"><p><sup>f</sup>MADRS: Montgomery and &#x00C5;sberg Depression Rating Scale.</p></fn><fn id="table1fn7"><p><sup>g</sup>YMRS: Young Mania Rating Scale</p></fn><fn id="table1fn8"><p><sup>h</sup>FAST: Functioning Assessment Short Test.</p></fn><fn id="table1fn9"><p><sup>i</sup>CGI-S: Clinical Global Impression Severity.</p></fn><fn id="table1fn10"><p><sup>j</sup>GAF: Global Assessment of Functioning.</p></fn><fn id="table1fn11"><p><sup>k</sup>FEBD: first episode bipolar disorder.</p></fn><fn id="table1fn12"><p><sup>l</sup>HR: hazard ratio.</p></fn><fn id="table1fn13"><p><sup>m</sup>AUC: area under the curve.</p></fn><fn id="table1fn14"><p><sup>n</sup>NLP: natural language processing.</p></fn><fn id="table1fn15"><p><sup>o</sup>ECE: expected calibration error.</p></fn><fn id="table1fn16"><p><sup>p</sup>ATE: average treatment effect.</p></fn><fn id="table1fn17"><p><sup>q</sup>ITR: individual treatment rule.</p></fn><fn id="table1fn18"><p><sup>r</sup>PPV: positive predictive value.</p></fn><fn id="table1fn19"><p><sup>s</sup>NPV: negative predictive value.</p></fn><fn id="table1fn20"><p><sup>t</sup>BrainAge indicates the difference between the estimated brain age and the chronological age expressed in years.</p></fn><fn id="table1fn21"><p><sup>u</sup>HAMD: Hamilton Depression Scale.</p></fn><fn id="table1fn22"><p><sup>v</sup>MSE: mean squared error.</p></fn><fn id="table1fn23"><p><sup>w</sup>AC: actual concentration [mmol/L].</p></fn><fn id="table1fn24"><p><sup>x</sup>PC: predictive concentration [mmol/L].</p></fn><fn id="table1fn25"><p><sup>y</sup>D%: deviation (%).</p></fn><fn id="table1fn26"><p><sup>z</sup>r: correlation coefficient.</p></fn><fn id="table1fn27"><p><sup>aa</sup>MAE: mean absolute error.</p></fn><fn id="table1fn28"><p><sup>ab</sup>MRE: mean relative error.</p></fn></table-wrap-foot></table-wrap></sec><sec id="s3-4"><title>Outcome Comparison</title><p>The target outcomes for AI models are reported below for the 5 main categories.</p><sec id="s3-4-1"><title>Acute Symptomatic Response</title><p>These studies [<xref ref-type="bibr" rid="ref35">35</xref>,<xref ref-type="bibr" rid="ref37">37</xref>,<xref ref-type="bibr" rid="ref38">38</xref>,<xref ref-type="bibr" rid="ref41">41</xref>,<xref ref-type="bibr" rid="ref46">46</xref>,<xref ref-type="bibr" rid="ref49">49</xref>,<xref ref-type="bibr" rid="ref59">59</xref>,<xref ref-type="bibr" rid="ref61">61</xref>] aimed to predict short-term symptom reduction within a 4/12-week timeframe from intervention. Pharmacological prediction models showed modest performance overall. For instance, a model developed for esketamine treatment achieved an accuracy of 68.5% [<xref ref-type="bibr" rid="ref49">49</xref>], while models predicting response to lithium and quetiapine based only on clinical variables explained 17.4% and 32.1% of the outcome variance, respectively [<xref ref-type="bibr" rid="ref37">37</xref>]. Imaging-based biomarkers seemed to provide improved predictive performance. Approaches based on structural connectivity matrices or low-frequency brain activity fluctuations reported accuracies ranging from 74% [<xref ref-type="bibr" rid="ref38">38</xref>] to 76.9% [<xref ref-type="bibr" rid="ref61">61</xref>]. However, not all neuroimaging approaches yielded positive results. Models based on functional magnetic resonance imaging (fMRI) data acquired during cognitive tasks alone did not reliably predict treatment status with erythropoietin, with accuracies not exceeding 60% [<xref ref-type="bibr" rid="ref46">46</xref>]. <xref ref-type="fig" rid="figure3">Figure 3</xref> summarizes the predictive performance of AI models for acute symptomatic response studies, showing a pooled AUC of 67.8% (95% CI 61.8%&#x2010;73.8%), an <italic>I</italic><sup>2</sup> equal to 93.6%, and a &#x03C4;<sup>2</sup> of 0.0074.</p><fig position="float" id="figure3"><label>Figure 3.</label><caption><p>Forest plot of AI model performance for acute symptomatic response prediction in bipolar disorder spectrum, expressed as area under the curve (AUC) with 95% CIs. Studies are displayed in descending order of AUC to facilitate visual comparison of model performance. Marker sizes are proportional to test-set sample size, and the corresponding study weights are reported on the right for transparent evaluation. The width of the diamond represents the 95% CI of the pooled AUC [<xref ref-type="bibr" rid="ref41">41</xref>,<xref ref-type="bibr" rid="ref46">46</xref>,<xref ref-type="bibr" rid="ref49">49</xref>,<xref ref-type="bibr" rid="ref61">61</xref>].</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="mental_v13i1e93307_fig03.png"/></fig></sec><sec id="s3-4-2"><title>Long-Term Maintenance Response</title><p>The primary aim of studies in this group [<xref ref-type="bibr" rid="ref32">32</xref>,<xref ref-type="bibr" rid="ref33">33</xref>,<xref ref-type="bibr" rid="ref36">36</xref>,<xref ref-type="bibr" rid="ref42">42</xref>,<xref ref-type="bibr" rid="ref43">43</xref>,<xref ref-type="bibr" rid="ref47">47</xref>,<xref ref-type="bibr" rid="ref53">53</xref>,<xref ref-type="bibr" rid="ref55">55</xref>-<xref ref-type="bibr" rid="ref57">57</xref>] was to predict long-term clinical stability (typically &#x003E;1&#x2010;2 years) provided by pharmacological interventions (mood stabilizers). Three main methodological approaches emerge. Large-scale clinical studies based on clinical interviews or EHR generally reported moderate but realistic predictive performance [<xref ref-type="bibr" rid="ref32">32</xref>,<xref ref-type="bibr" rid="ref36">36</xref>,<xref ref-type="bibr" rid="ref47">47</xref>]. In contrast, biological and cellular models using biomarkers such as induced pluripotent stem cells or gene expression data achieved exceptionally high accuracies, ranging from 95.8% [<xref ref-type="bibr" rid="ref43">43</xref>] to 99% [<xref ref-type="bibr" rid="ref43">43</xref>,<xref ref-type="bibr" rid="ref57">57</xref>]. However, these findings were typically based on relatively small sample sizes, which may limit their generalizability. Genetic and epigenetic approaches, including DNA methylation&#x2013;based models and polygenic risk scores, showed intermediate predictive performance, with AUC values ranging from 70 [<xref ref-type="bibr" rid="ref42">42</xref>] to 86% [<xref ref-type="bibr" rid="ref53">53</xref>]. Notably, these models tended to demonstrate greater accuracy in identifying nonresponders than responders. <xref ref-type="fig" rid="figure4">Figure 4</xref> summarizes the predictive performance of AI models for long-term maintenance response studies&#x2014;specifically, those referring to lithium response&#x2014;showing a pooled AUC of 80.2% (95% CI 74.3%&#x2010;86.1%), an <italic>I</italic><sup>2</sup> of 56.3%, and a &#x03C4;<sup>2</sup> of 0.0022.</p><fig position="float" id="figure4"><label>Figure 4.</label><caption><p>Forest plot of AI model performance for long-term maintenance response prediction in the bipolar disorder spectrum, expressed as area under the curve (AUC) with 95% CIs. Studies are displayed in descending order of AUC to facilitate visual comparison of model performance. Marker sizes are proportional to test-set sample size, and the corresponding study weights are reported on the right for transparent evaluation. The width of the diamond represents the 95% CI of the pooled AUC [<xref ref-type="bibr" rid="ref33">33</xref>,<xref ref-type="bibr" rid="ref42">42</xref>,<xref ref-type="bibr" rid="ref47">47</xref>].</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="mental_v13i1e93307_fig04.png"/></fig></sec><sec id="s3-4-3"><title>Relapse and Readmission Risk</title><p>The focus of studies from this group [<xref ref-type="bibr" rid="ref22">22</xref>,<xref ref-type="bibr" rid="ref34">34</xref>,<xref ref-type="bibr" rid="ref39">39</xref>,<xref ref-type="bibr" rid="ref40">40</xref>,<xref ref-type="bibr" rid="ref45">45</xref>,<xref ref-type="bibr" rid="ref48">48</xref>,<xref ref-type="bibr" rid="ref51">51</xref>,<xref ref-type="bibr" rid="ref54">54</xref>] was on predicting critical outcomes, such as rehospitalization and behavioral relapse. Large-scale studies based on national registries and big data provided robust and generalizable estimates. For example, the studies conducted on Swedish and Finnish cohorts (N&#x2248;45,000) [<xref ref-type="bibr" rid="ref39">39</xref>,<xref ref-type="bibr" rid="ref40">40</xref>], reported stable AUC values ranging from 0.68 to 0.77 for predicting mortality and 2-year relapse. Similarly, the random forest (RF) models showed that rehospitalization was associated with a significant reduction in suicide attempts, but only within specific high-risk subgroups [<xref ref-type="bibr" rid="ref51">51</xref>]. More recent approaches based on digital phenotyping also showed promising results. The use of wearable devices enabled the differentiation between depressive and manic episodes with AUC values between 0.85 and 0.88 [<xref ref-type="bibr" rid="ref22">22</xref>]. Finally, AI rule&#x2013;based approaches have also been explored. Inductive logic programming to identify relapse patterns achieved an accuracy of 85% [<xref ref-type="bibr" rid="ref54">54</xref>]. <xref ref-type="fig" rid="figure5">Figure 5</xref> summarizes the predictive performance of AI models for relapse and readmission risk studies in a forest plot, showing a pooled AUC of 71.2% (95% CI 68.8%&#x2010;73.7%), an <italic>I</italic><sup>2</sup> equals to 93.0%, and a &#x03C4;<sup>2</sup> of 0.0039.</p><fig position="float" id="figure5"><label>Figure 5.</label><caption><p>Forest plot of AI model performance for relapse and readmission risk in the bipolar disorder spectrum, expressed as AUC with 95% CIs. Studies are displayed in descending order of AUC to facilitate visual comparison of model performance. Marker sizes are proportional to test-set sample size, and the corresponding study weights are reported on the right for transparent evaluation. The width of the diamond represents the 95% CI of the pooled AUC [<xref ref-type="bibr" rid="ref34">34</xref>,<xref ref-type="bibr" rid="ref39">39</xref>,<xref ref-type="bibr" rid="ref40">40</xref>,<xref ref-type="bibr" rid="ref48">48</xref>].</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="mental_v13i1e93307_fig05.png"/></fig></sec><sec id="s3-4-4"><title>Safety and Dose Optimization</title><p>In this group, the main outcomes are no longer symptomatic responses but safety and pharmacokinetic measures. Several studies used ML to support dose optimization, with CatBoost and neural network models reaching about 85% accuracy for valproate dosing [<xref ref-type="bibr" rid="ref62">62</xref>] and a 97.3% [<xref ref-type="bibr" rid="ref60">60</xref>] correlation between predicted and observed lithium blood concentrations. Other models were used to predict adverse events, including prolactin levels, with a relative error of 11% [<xref ref-type="bibr" rid="ref63">63</xref>]. Beyond predictive accuracy, some studies also used target maximum likelihood estimation and Super Learner ML-based statistical algorithms [<xref ref-type="bibr" rid="ref30">30</xref>,<xref ref-type="bibr" rid="ref50">50</xref>] to estimate average treatment effects in real-world data, identifying small but potentially relevant absolute differences between treatments.</p></sec><sec id="s3-4-5"><title>Brain Aging and Phenotyping</title><p>In this group, the outcome of interest is defined as biological deviation from normative trajectories or the identification of novel diagnostic subgroups. Two studies examined brain aging as a transdiagnostic biomarker [<xref ref-type="bibr" rid="ref52">52</xref>,<xref ref-type="bibr" rid="ref58">58</xref>]. Deviations from expected aging patterns were quantified using metrics such as the BrainAge index and the Quantile Regression Index (QRI). Findings suggest that BD is associated with accelerated brain aging, with a moderate effect size (Cohen <italic>d</italic>=0.55) [<xref ref-type="bibr" rid="ref52">52</xref>]. Notably, lithium treatment appears to have a protective effect, with patients showing brain age estimates more closely resembling those of healthy controls [<xref ref-type="bibr" rid="ref58">58</xref>]. Other approaches aimed to redefine clinical heterogeneity through data-driven clustering. Four treatment-based clusters were found to better explain longitudinal functional outcomes, as measured by Global Assessment of Functioning (GAF) and Functioning Assessment Short Test (FAST) scores, compared to traditional <italic>DSM (Diagnostic and Statistical Manual of Mental Disorders)</italic>&#x2013;based diagnostic categories, as reflected by improved statistical model fit indices [<xref ref-type="bibr" rid="ref31">31</xref>].</p></sec></sec><sec id="s3-5"><title>Cross-Domain Overview</title><p><xref ref-type="fig" rid="figure6">Figure 6</xref> provides a comparative overview of AI model performance across different marker domains (ie, clinical, imaging, genomics, and digital-wearable). For studies assessing multiple models, only the main models are shown, either the best-performing models or the ones developed specifically for the study. Clinical-based models were represented across the full range of models, achieving better accuracies, although the highest average accuracies were observed in imaging- and genomics-based models, which were often derived from comparatively smaller BD sample sizes (smaller bubbles).</p><fig position="float" id="figure6"><label>Figure 6.</label><caption><p>Bubble chart showing the performance of different AI models across various domains. The data used for the chart creation are reported in Table S2 in the <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>. SVM: support vector machine.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="mental_v13i1e93307_fig06.png"/></fig></sec><sec id="s3-6"><title>Bias Assessment Risk</title><p>The potential publication bias showed that studies with small sample sizes tended to report exceptionally high AUC values, often close to perfect, while as the sample size increased, the AUC values converged toward a more moderate and realistic range (<xref ref-type="fig" rid="figure7">Figure 7</xref>). The Egger test did not reveal evidence of funnel plot asymmetry (intercept=&#x2212;4.63; <italic>P</italic>=.56).</p><fig position="float" id="figure7"><label>Figure 7.</label><caption><p>Funnel plot assessing potential publication bias, with area under the curve (AUC) on the X-axis and the inverse square root of sample size on the Y-axis, illustrating the relationship between model performance and study size; points in the lower part of the plot indicate smaller sample sizes, while points further to the right indicate higher AUC values. The gray shaded area denotes the infeasible region where AUC values are undefined. For each study, only the model performing with the best AUC is shown. AUC: area under the curve.</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="mental_v13i1e93307_fig07.png"/></fig><p>The methodological evaluation using PROBAST+AI found that the great majority of the included studies [<xref ref-type="bibr" rid="ref22">22</xref>,<xref ref-type="bibr" rid="ref30">30</xref>-<xref ref-type="bibr" rid="ref63">63</xref>] were at high risk of bias (<xref ref-type="fig" rid="figure8">Figure 8</xref>). In the model development phase, most studies were judged at high risk, primarily due to limitations in the analysis domain (Figure S1 in <xref ref-type="supplementary-material" rid="app1">Multimedia Appendix 1</xref>), including inadequate handling of overfitting (ie, a model that performs better on training data than on independent data), insufficient reporting of model performance measures, and lack of external validation. While the participants&#x2019; and predictors&#x2019; domains were frequently rated as low risk, concerns regarding predictors&#x2019; applicability and overall methodological quality contributed to unfavorable overall judgments.</p><p>Similarly, in the model evaluation phase, the overall risk of bias was predominantly rated as high. This was mainly due to shortcomings in the analysis domain, such as limited external validation, inappropriate performance assessment, or insufficient sample sizes. Although predictors and outcomes were often judged as low or unclear risk, these did not offset the high risk associated with analytical issues. Overall applicability in the evaluation phase was judged as low in a small subset of studies but remained unclear or high in most cases, reflecting heterogeneity in validation settings with limited generalizability and fairness.</p><fig position="float" id="figure8"><label>Figure 8.</label><caption><p>Traffic-light plot of overall judgments for the assessed studies using the Prediction model Risk Of Bias Assessment Tool for prediction models using regression or AI methods (PROBAST+AI). Panel A shows the overall judgments for model development. Panel B shows the overall judgments for model evaluation. Colors indicate the level of risk: green=low, yellow=unclear, red=high. Each row represents a study, and each column corresponds to a specific overall judgment domain [<xref ref-type="bibr" rid="ref22">22</xref>,<xref ref-type="bibr" rid="ref30">30</xref>-<xref ref-type="bibr" rid="ref63">63</xref>].</p></caption><graphic alt-version="no" mimetype="image" position="float" xlink:type="simple" xlink:href="mental_v13i1e93307_fig08.png"/></fig></sec></sec><sec id="s4" sec-type="discussion"><title>Discussion</title><sec id="s4-1"><title>Principal Findings</title><p>The present systematic review synthesized evidence from 35 studies investigating the application of AI methodologies to treatment optimization in BD, restricting inclusion to studies published from 2015 onward, in line with the marked growth of research on classical AI in psychiatry during this period [<xref ref-type="bibr" rid="ref71">71</xref>]. Across the literature, AI approaches were primarily used for five clinical purposes: (1) acute symptomatic response, (2) long-term maintenance response, (3) relapse and readmission risk, (4) safety and dose optimization, and (5) brain aging and phenotyping. While the rapid growth of studies reflects increasing interest in precision psychiatry from 2020 onward, our critical analysis revealed substantial methodological weaknesses that limit the clinical utility of most models. Although the pooled performance metrics across the 5 categories appear encouraging at first glance, the high heterogeneity across studies and the predominance of high risk of bias identified by the PROBAST+AI assessment suggest that these results must be handled with caution. In most cases, reported accuracies represent optimistic estimates derived from small or internally validated samples rather than robust, generalizable clinical tools.</p></sec><sec id="s4-2"><title>Acute Symptomatic Response</title><p>The forest plot for this group showed a heterogeneous pattern of results, with a modest pooled AUC indicating only moderate discriminative performance. This suggests that acute symptomatic response remains difficult to predict reliably, likely because these are short-term treatment outcomes influenced by multiple interacting clinical and biological factors. Due to the dynamic and multifactorial nature of acute treatment response, the high heterogeneity (<italic>I</italic>&#x00B2;=93.6%) reflects substantial differences that are unlikely to be explained by sampling error alone. Therefore, the pooled AUC estimate should be interpreted prudently, as it summarizes studies that differ markedly in treatment modality (eg, first or second generation of antidepressants), predictor type (clinical vs neuroimaging, with this one representing the most informative one), and methodological approaches.</p><p>Despite this, the estimated between-study variance was low in absolute terms (&#x03C4;&#x00B2;=0.0074). However, this does not eliminate concerns about heterogeneity. Rather, it suggests that the dispersion of effects on the AUC was limited in absolute magnitude, but variability within studies is large. This combination is consistent with a modest AUC that was not consistently reproduced across the studies, particularly in neuroimaging-based models.</p><p>Differences in validation strategies, with more rigorous cross-site validation yielding lower but more realistic performance, further contribute to the observed variability. Overall, these findings suggest that AI models may detect a biological signal associated with acute response, but their current predictive utility is limited and generalizability remains uncertain.</p></sec><sec id="s4-3"><title>Long-Term Maintenance Response</title><p>This group showed clinically meaningful predictive performance, indicating fair discriminative ability of AI models in identifying pharmacological (lithium) responders. Between-study heterogeneity was moderate (<italic>I</italic>&#x00B2;=56.3%), suggesting that some of the variability across study results reflected genuine differences. Moreover, the estimated between-study variance was low in absolute terms (&#x03C4;&#x00B2;=0.0022), suggesting limited dispersion of the underlying effects on the AUC. These findings highlight that, although methodological and clinical differences across studies represent a constraint, they do not preclude a consistent overall predictive performance comparison and reinforce the idea that drug-treatment response (lithium) is a predictable target. However, the observed heterogeneity suggests that improvements in predictive accuracy may depend on the multimodal integration of sociodemographic, clinical, and biological data, in line with the current trend in precision psychiatry, with genetic predictors appearing the most informative. Additionally, assessment using PROBAST+AI identified a high risk of bias in each domain across almost all studies in this group; therefore, the findings must be interpreted cautiously. This category included 4 studies characterized by significant publication bias due to small sample size associated with very high AUC [<xref ref-type="bibr" rid="ref43">43</xref>,<xref ref-type="bibr" rid="ref53">53</xref>,<xref ref-type="bibr" rid="ref57">57</xref>] or performance outside the confidence interval [<xref ref-type="bibr" rid="ref56">56</xref>], which must be minimized in future AI development efforts.</p></sec><sec id="s4-4"><title>Relapse and Readmission Risk</title><p>In this group, AI models achieved moderate predictive performance for clinically critical outcomes, indicating fair-to-good discriminative ability for identifying patients at risk of relapse, hospitalization, or mortality. This performance is supported by both high-accuracy models in specific contexts and the stability of large registry-based studies. The most informative predictors were wearable and clinical ones. Interpretation of this pooled estimate requires caution because between-study heterogeneity was high (<italic>I</italic>&#x00B2;=93.0%), reflecting differences in outcome definitions (eg, manic relapse, suicidal behavior, and all-cause mortality), data modalities, and sample size, which ranged from large-scale EHR-based studies to smaller multimodal approaches integrating neuroimaging. Despite this variability, the low between-study variance (&#x03C4;&#x00B2;=0.0039) suggests that the pooled AUC may still capture a broad average predictive signal consistent across studies, indicating a shared set of clinically relevant risk factors. Indeed, converging evidence highlights the importance of variables such as prior self-harm, hospitalization history, and medical comorbidities. Moreover, despite the differences in sample size, the weights assigned to each model in the forest plot are similar, indicating that no single study disproportionately influences the pooled estimate of predictive performance. From a clinical perspective, these models showed potential to support proactive risk stratification and personalized monitoring strategies, with external validations across different health care systems further supporting their generalizability. Overall, although methodological fragmentation persists, the findings indicate that AI can reliably capture the trajectory of illness severity in BD. Stronger conclusions about generalizability will require more homogeneous outcome definitions and more consistent external validation.</p></sec><sec id="s4-5"><title>Safety and Dose Optimization</title><p>The group of studies focusing on safety and dose optimization highlights the translational potential of AI in moving beyond generalized clinical guidelines toward individualized treatment strategies based on patient-specific risk and pharmacokinetic profiles. AI models showed strong performance in predicting blood concentrations of key mood stabilizers, with approaches such as artificial neural networks and gradient boosting algorithms identifying clinically relevant predictors (ie, age, renal function, co-medications, and bilirubin levels) to support personalized dosing. In parallel, advanced causal inference methods have been applied to better quantify the cardiometabolic risks associated with antipsychotic treatments, enabling more precise risk stratification and informing safer treatment switches using informative predictors, such as age, sex, baseline health status, psychiatric comorbidities, previous exposure to medications with cardiometabolic effects, and other clinical data. AI-driven pharmacovigilance has also shown promise in identifying predictors of adverse effects, including endocrine complications such as antipsychotic-induced hyperprolactinemia. Furthermore, the integration of NLP applied to EHR has revealed important discrepancies between real-world prescribing patterns and controlled trial evidence, underscoring the need for AI-based clinical decision support systems. These findings suggest that AI can play a key role in optimizing both the efficacy and safety of pharmacological treatments in BD, facilitating a more precise and data-driven approach to clinical management.</p></sec><sec id="s4-6"><title>Brain Aging and Phenotyping</title><p>Studies focusing on phenotyping and brain aging highlight the potential of AI to move beyond traditional diagnostic categories by identifying biologically and clinically meaningful dimensions of disease progression. Unsupervised ML approaches indicate that clusters based on real-world treatment patterns explain longitudinal functional outcomes better than the conventional BD-I/BD-II distinction, advising that illness trajectories are closely linked to specific pharmacological response phenotypes. In parallel, AI-derived metrics such as BrainAGE and the QRI have revealed, through informative imaging predictors, an accelerated brain aging process in BD, particularly affecting white matter and subcortical structures, driven independently by illness severity and cardiometabolic comorbidities, with hypertension emerging as a key risk factor. Notably, converging evidence supports a neuroprotective role for lithium, which is associated with a younger brain age comparable to that of healthy controls. Overall, these findings suggest that AI can provide objective biomarkers of &#x201C;brain frailty,&#x201D; paving the way for precision psychiatry approaches aimed at monitoring disease progression and personalizing interventions to mitigate long-term cognitive decline. However, the limited number of studies in this group precludes definitive conclusions.</p></sec><sec id="s4-7"><title>Comprehensive Overview</title><p>The relationship between AI models and classification accuracy across different application domains in BD research highlights substantial heterogeneity in both performance and data typology. Gradient boosting models were most frequently adopted, particularly in the clinical domain, where they achieve a wide range of accuracies (from random to good) and were often trained on larger sample sizes due to the relative ease of collecting such data. DL models tended to achieve higher accuracies (up to approximately 0.95), although typically with more restricted sample sizes, raising potential concerns regarding their downstream generalizability. Genomics-based studies, despite smaller cohorts, frequently reported high accuracies even with different models, suggesting strong informative signal-to-noise ratios but with the highest risks of overfitting. Wearable-based applications remain limited in number, although they demonstrate promise. This emphasizes the need for balanced methodological choices during model development to ensure robust, unbiased, and clinically applicable AI systems for patients with BD.</p><p>Although the Egger test did not reveal evidence of funnel plot asymmetry, its interpretation is limited by the small number of studies reporting AUC. Visual inspection of the funnel plot suggests that most small-sample studies are distributed on the right side of the central dashed line (AUC&#x2248;0.8), indicating estimates of high performance. Conversely, there is a clear gap in the lower left portion, where small studies with poor performance should be located. This distribution suggests that scientific journals tend to publish small studies only when the results are very good, leading to a systematic overestimation of the true effectiveness of AI. Studies with modest performance but small samples probably remain unpublished or confidential. As the sample size increases, AUC values tend to converge toward a more realistic range. Attention should be moved from &#x201C;near-perfect&#x201D; AUC in small samples to more robust, externally validated models that demonstrate stable and reliable performance in heterogeneous populations.</p><p>Across the 5 categories, PROBAST+AI analysis revealed that most studies were at high risk of bias, particularly in the analysis domain. Indeed, only 2 [<xref ref-type="bibr" rid="ref37">37</xref>,<xref ref-type="bibr" rid="ref40">40</xref>] studies received a low-risk judgment across all domains of the PROBAST+AI assessment, due to rigorous data handling, robust validation (including external validation), adherence to standards, and comprehensive evaluation of model performance, including calibration and fairness. In the other studies, applicability concerns were common, reflecting nontransparent AI system modeling or limited generalizability to real-world clinical settings. Only a minority of studies demonstrated consistently low risk across domains, while several judgments remain unclear due to insufficient methodological reporting. Furthermore, only 7 [<xref ref-type="bibr" rid="ref36">36</xref>,<xref ref-type="bibr" rid="ref39">39</xref>,<xref ref-type="bibr" rid="ref40">40</xref>,<xref ref-type="bibr" rid="ref43">43</xref>,<xref ref-type="bibr" rid="ref45">45</xref>,<xref ref-type="bibr" rid="ref48">48</xref>,<xref ref-type="bibr" rid="ref50">50</xref>] out of the 35 included studies provided access to their models, underscoring a lack of transparency and availability of predictive tools, which limits reproducibility and the possibility of external validation. In light of the limitations identified through the PROBAST+AI assessment, future model development in the context of BD spectrum treatment should focus on enhancing dataset diversity to reduce bias, improving external validation across different clinical settings, and integrating clinically interpretable features to support decision-making. Additionally, incorporating longitudinal patient data and multimodal inputs may strengthen predictive performance and generalizability, ensuring that AI-driven recommendations are both robust and clinically applicable.</p><p>Data-related challenges represented a major constraint. Multicenter datasets were often affected by differences in clinical protocols and data acquisition procedures, whereas high-dimensional modalities such as omics introduced technical artifacts, including batch effects, that can distort predictive signals. This variability undermines AI models&#x2019; stability and reproducibility across independent cohorts. A further concern was the imbalance between the number of predictors and the available sample sizes. Many studies [<xref ref-type="bibr" rid="ref43">43</xref>,<xref ref-type="bibr" rid="ref57">57</xref>] relied on complex omics or neuroimaging features derived from very small cohorts, creating conditions prone to overfitting. Although some studies [<xref ref-type="bibr" rid="ref37">37</xref>,<xref ref-type="bibr" rid="ref41">41</xref>,<xref ref-type="bibr" rid="ref43">43</xref>] used regularization or feature selection, these strategies were not externally validated. Calibration and real-world integration were rarely assessed. Most models were evaluated primarily using discrimination metrics, leaving their practical impact in clinical scenarios uncertain.</p><p>From a regulatory perspective, all psychiatric AI systems are likely to be classified as high-risk under the EU AI-Act [<xref ref-type="bibr" rid="ref72">72</xref>], requiring even stronger evidence of robustness, technical documentation, transparency, fairness, and ethically grounded development.</p><p>Several regulatory and methodological frameworks have been proposed to guide the responsible development and application of both classical and generative AI in the life sciences and health care sectors. The European Medicines Agency (EMA) has issued a discussion document on the use of AI throughout the medicinal product lifecycle, outlining considerations for the safe and effective integration of AI/ML at all stages from drug development to postauthorization, while emphasizing the need to align its use with EU regulatory requirements and data governance principles [<xref ref-type="bibr" rid="ref73">73</xref>]. In addition, EMA guidance on AI highlights the importance of ensuring transparency, robustness, and appropriate validation of AI-based tools to support regulatory decision-making. Internationally, initiatives such as the Food and Drug Administration (FDA) draft guidance on AI-enabled medical products [<xref ref-type="bibr" rid="ref74">74</xref>] and the joint Good Machine Learning Practice (GMLP) guiding principles developed by the FDA, Health Canada, and the Medicines and Healthcare Products Regulatory Agency (MHRA) [<xref ref-type="bibr" rid="ref75">75</xref>] further emphasize life cycle&#x2013;wide considerations, including data quality, model performance monitoring, and documentation to ensure safety and effectiveness in real-world applications. Furthermore, consensus frameworks such as Fair Universal Traceable Usable Robust Explainable&#x2013;AI (FUTURE-AI) [<xref ref-type="bibr" rid="ref76">76</xref>] propose structured best practices covering the entire AI lifecycle in health care, from design and development to deployment and monitoring, with an emphasis on trustworthiness, explainability, and clinical validity. Collectively, these methodological recommendations emphasize the importance of a risk-based, transparent, and continuously evaluated approach to AI implementation in the life sciences.</p><p>The medico-legal status of AI in the life sciences remains complex and continues to evolve. Although AI-based tools are increasingly used to support research, drug development, and clinical decision-making, their legal framing depends on their intended use and degree of autonomy. In the EU, AI systems fall under the regulatory framework for medical devices when used for diagnostic or therapeutic purposes, requiring compliance with the Medical Device Regulation (EU 2017/745 MDR) [<xref ref-type="bibr" rid="ref77">77</xref>] and, where applicable, the In Vitro Diagnostic Regulation (EU 2017/746 IVDR) [<xref ref-type="bibr" rid="ref78">78</xref>]. As such, the safe integration of AI into life sciences requires not only technical validation but also regulatory compliance and robust governance frameworks.</p><p>Several systematic reviews have highlighted the increased use of AI models in psychiatric disorders beyond BD, particularly in major depressive disorder, schizophrenia, anxiety disorders, and suicide risk prediction. In major depressive disorder, ML approaches applied to clinical, neuroimaging, and digital phenotyping data have shown moderate to high performance in identifying diagnoses and predicting treatment response, although generalizability remains limited by small sample sizes and heterogeneous study designs [<xref ref-type="bibr" rid="ref79">79</xref>,<xref ref-type="bibr" rid="ref80">80</xref>]. In schizophrenia, AI models have been used to support diagnosis and symptom classification using neuroimaging and speech-based features, demonstrating promising accuracy but requiring further external validation [<xref ref-type="bibr" rid="ref14">14</xref>]. Similarly, in anxiety disorders, wearable and smartphone-derived data have been explored to detect symptom fluctuations and stress-related patterns, although methodological variability and lack of standardized endpoints limit clinical translation [<xref ref-type="bibr" rid="ref81">81</xref>]. A substantial body of literature has also focused on suicide risk prediction, where NLP and multimodal ML models have achieved encouraging predictive performance using EHR and social media data, but raise important concerns regarding interpretability, ethical use, and false-positive rates in clinical practice [<xref ref-type="bibr" rid="ref82">82</xref>,<xref ref-type="bibr" rid="ref83">83</xref>].</p><p>The current evidence suggests that, consistently with other psychiatric condition models, AI in BD treatment is still largely in its exploratory or proof-of-concept phase. While the conceptual shift toward precision psychiatry is evident, most existing models are not yet suitable for routine clinical decision-making.</p><p>The translational gap does not lie in algorithmic sophistication but rather in: (1) data quality and representativeness; (2) validation strategies; (3) transparent and interpretable modeling from human experts (ie, human in the loop); and (4) prospective clinical evaluation. Until these issues are addressed, AI tools remain prone to the risk of producing optimistic results at the benchside that are not at the bedside and in real-world settings.</p><p>Overall, although AI applications in psychiatry show significant potential across multiple conditions, the current evidence remains largely exploratory, and robust prospective validation studies are needed before real-world clinical implementation.</p></sec><sec id="s4-8"><title>Limitations</title><p>Several limitations should be considered when interpreting the findings of this review. The included studies showed heterogeneity in sample characteristics, outcome definitions, follow-up duration, relapse types, and treatment strategies, which was only partially mitigated using broad search strategies. This variability reduced comparability and limited the interpretability of pooled performance estimates. Furthermore, many investigations were based on small cohorts or exploratory designs, restricting the generalizability of their findings.</p><p>The conclusions of this review should be interpreted in light of several methodological limitations inherent in this literature. First, there is a lack of standardized diagnostic criteria for BD that use heterogeneous definitions across <italic>ICD-9/10</italic> (<italic>International Classification of Diseases, 9/10th Revision)</italic>, <italic>DSM-IV (Diagnostic and Statistical Manual of Mental Disorders</italic> [Fourth Edition]), and <italic>DSM-5</italic> (<italic>Diagnostic and Statistical Manual of Mental Disorders</italic> [Fifth Edition]), complicating direct comparisons and limiting generalizability. A further limitation is that we did not account for time since diagnosis or the presence of complex medical and psychiatric comorbidities, often relying instead on static factors that fail to capture the longitudinal progression of the disease. Moreover, our analysis was restricted to peer-reviewed publications, excluding proprietary or industrial models protected by patents or trade secrecy, which may perform differently but are not accessible for independent scientific validation. Another key limitation of this review is that none of the included studies addressed nonpharmacological interventions. This restricts the generalizability of our findings and prevents a comprehensive assessment of AI tools across the full spectrum of treatment strategies. Although this review focuses on AI applications, most models rely on ML techniques, with other classical AI approaches being underrepresented. This limited diversity of methods should be considered a potential limitation of the current evidence. Finally, this systematic review was not prospectively registered. While preregistration is considered best practice to enhance transparency and minimize the risk of selective outcome reporting, the review was conducted in accordance with PRISMA guidelines, and all methodological steps, including eligibility criteria and outcomes, were defined a priori and consistently applied. Nonetheless, the absence of registration remains a limitation.</p></sec><sec id="s4-9"><title>Future Directions</title><p>Future research should prioritize methodological maturity and translationality over incremental performance gains. This will require larger, harmonized multicenter benchmarking datasets on which to perform systematic external validation and to favor AI-model calibration, real decision impact, and workflow integration. The development of open, shared repositories (eg, ENIGMA [Enhancing NeuroImaging Genetics through Meta-Analysis; ENIGMA Consortium] and PsychENCODE [Psychiatric Encyclopedia of DNA Elements; PsychENCODE Consortium]) and public computational e-infrastructures (eg, OpenNeuro [Stanford Center for Reproducible Neuroscience], Zenodo [CERN], NewPsy4U [Laboratory of Neuroinformatics, IRCCS Istituto Centro San Giovanni di Dio Fatebenefratelli]) will be propaedeutic to improve reproducibility and support the transition from experimental models to clinically reliable AI systems in BD.</p></sec><sec id="s4-10"><title>Conclusion</title><p>Classical AI approaches show potential for improving treatment response prediction, relapse risk estimation, and patient stratification in BD, supporting the transition toward precision psychiatry. However, the current evidence base remains methodologically fragile. Most studies showed a high risk of bias, limited external validation, small and heterogeneous samples, and inconsistent outcome definitions when assessed with the PROBAST+AI method. Consequently, reported performance metrics are likely to be overly optimistic and, at present, not generalizable to routine clinical practice.</p><p>Progress will depend on rigorous study design, standardized methods, including benchmarking datasets, transparent reporting, and prospective validation to bridge the gap between proof-of-concept studies and deployable clinical decision support systems.</p></sec></sec></body><back><ack><p>The authors disclose the use of generative AI tools (ChatGPT by OpenAI, NotebookLM by Google, and Perplexity AI) to support language editing and improvement of clarity of the manuscript. All scientific content, interpretations, and conclusions were critically reviewed and validated by the authors, who take full responsibility for the final version of the manuscript.</p></ack><notes><sec><title>Funding</title><p>The present work was supported by &#x201C;Ministero della Salute&#x201D;, IRCCS Research Program, Ricerca Corrente - Linea n. 1 &#x201C;Utilizzo di strumenti di Intelligenza Artificiale (AI) per l&#x2019;analisi dei disturbi psichici&#x201D;.</p></sec></notes><fn-group><fn fn-type="conflict"><p>None declared.</p></fn></fn-group><glossary><title>Abbreviations</title><def-list><def-item><term id="abb1">AUC</term><def><p>area under the curve</p></def></def-item><def-item><term id="abb2">BD</term><def><p>bipolar disorder</p></def></def-item><def-item><term id="abb3">CNN</term><def><p>convolutional neural network</p></def></def-item><def-item><term id="abb4">DL</term><def><p>deep learning</p></def></def-item><def-item><term id="abb5"><italic>DSM-5</italic></term><def><p><italic>Diagnostic and Statistical Manual of Mental Disorders</italic> (Fifth Edition)</p></def></def-item><def-item><term id="abb6"><italic>DSM-IV</italic></term><def><p><italic>Diagnostic and Statistical Manual of Mental Disorders</italic> (Fourth Edition)</p></def></def-item><def-item><term id="abb7">EHR</term><def><p>electronic health record</p></def></def-item><def-item><term id="abb8">EMA</term><def><p>European Medicines Agency</p></def></def-item><def-item><term id="abb9">ENIGMA</term><def><p>Enhancing NeuroImaging Genetics through Meta-Analysis</p></def></def-item><def-item><term id="abb10">EU</term><def><p>European Union</p></def></def-item><def-item><term id="abb11">FAST</term><def><p>Functioning Assessment Short Test</p></def></def-item><def-item><term id="abb12">FDA</term><def><p>Food and Drug Administration</p></def></def-item><def-item><term id="abb13">fMRI</term><def><p>functional magnetic resonance imaging</p></def></def-item><def-item><term id="abb14">FUTURE-AI</term><def><p>Fair Universal Traceable Usable Robust Explainable&#x2013;AI</p></def></def-item><def-item><term id="abb15">GAF</term><def><p>Global Assessment of Functioning</p></def></def-item><def-item><term id="abb16">GMLP</term><def><p>Good Machine Learning Practice</p></def></def-item><def-item><term id="abb17">HAMD</term><def><p>Hamilton Depression Scale</p></def></def-item><def-item><term id="abb18"><italic>ICD-9/10</italic></term><def><p><italic>International Classification of Diseases, 9/10th Revision</italic></p></def></def-item><def-item><term id="abb19">IVDR</term><def><p>In Vitro Diagnostic Regulation</p></def></def-item><def-item><term id="abb20">MDR</term><def><p>Medical Device Regulation</p></def></def-item><def-item><term id="abb21">MHRA</term><def><p>Medicines and Healthcare products Regulatory Agency</p></def></def-item><def-item><term id="abb22">ML</term><def><p>machine learning</p></def></def-item><def-item><term id="abb23">NLP</term><def><p>natural language processing</p></def></def-item><def-item><term id="abb24">PRISMA</term><def><p>Preferred Reporting Items for Systematic Reviews and Meta-Analyses</p></def></def-item><def-item><term id="abb25">PROBAST+AI</term><def><p>Prediction model Risk Of Bias Assessment Tool for prediction models using regression or AI methods</p></def></def-item><def-item><term id="abb26">PsychENCODE</term><def><p>Psychiatric Encyclopedia of DNA Elements</p></def></def-item><def-item><term id="abb27">QRI</term><def><p>Quantile Regression Index</p></def></def-item><def-item><term id="abb28">RF</term><def><p>random forest</p></def></def-item><def-item><term id="abb29">RNN</term><def><p>recurrent neural network</p></def></def-item><def-item><term id="abb30">SVM</term><def><p>support vector machine</p></def></def-item><def-item><term id="abb31">YMRS</term><def><p>Young Mania Rating Scale</p></def></def-item></def-list></glossary><ref-list><title>References</title><ref id="ref1"><label>1</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Burdick</surname><given-names>KE</given-names> </name><name name-style="western"><surname>Millett</surname><given-names>CE</given-names> </name><name name-style="western"><surname>Yocum</surname><given-names>AK</given-names> </name><etal/></person-group><article-title>Predictors of functional impairment in bipolar disorder: results from 13 cohorts from seven countries by the global bipolar cohort collaborative</article-title><source>Bipolar Disord</source><year>2022</year><month>11</month><volume>24</volume><issue>7</issue><fpage>709</fpage><lpage>719</lpage><pub-id pub-id-type="doi">10.1111/bdi.13208</pub-id><pub-id pub-id-type="medline">35322518</pub-id></nlm-citation></ref><ref id="ref2"><label>2</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Merikangas</surname><given-names>KR</given-names> </name><name name-style="western"><surname>Jin</surname><given-names>R</given-names> </name><name name-style="western"><surname>He</surname><given-names>JP</given-names> </name><etal/></person-group><article-title>Prevalence and correlates of bipolar spectrum disorder in the world mental health survey initiative</article-title><source>Arch Gen Psychiatry</source><year>2011</year><month>03</month><volume>68</volume><issue>3</issue><fpage>241</fpage><lpage>251</lpage><pub-id pub-id-type="doi">10.1001/archgenpsychiatry.2011.12</pub-id><pub-id pub-id-type="medline">21383262</pub-id></nlm-citation></ref><ref id="ref3"><label>3</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Grande</surname><given-names>I</given-names> </name><name name-style="western"><surname>Berk</surname><given-names>M</given-names> </name><name name-style="western"><surname>Birmaher</surname><given-names>B</given-names> </name><name name-style="western"><surname>Vieta</surname><given-names>E</given-names> </name></person-group><article-title>Bipolar disorder</article-title><source>Lancet</source><year>2016</year><month>04</month><day>9</day><volume>387</volume><issue>10027</issue><fpage>1561</fpage><lpage>1572</lpage><pub-id pub-id-type="doi">10.1016/S0140-6736(15)00241-X</pub-id><pub-id pub-id-type="medline">26388529</pub-id></nlm-citation></ref><ref id="ref4"><label>4</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Oliva</surname><given-names>V</given-names> </name><name name-style="western"><surname>Fico</surname><given-names>G</given-names> </name><name name-style="western"><surname>De Prisco</surname><given-names>M</given-names> </name><name name-style="western"><surname>Gonda</surname><given-names>X</given-names> </name><name name-style="western"><surname>Rosa</surname><given-names>AR</given-names> </name><name name-style="western"><surname>Vieta</surname><given-names>E</given-names> </name></person-group><article-title>Bipolar disorders: an update on critical aspects</article-title><source>Lancet Reg Health Eur</source><year>2025</year><month>01</month><volume>48</volume><fpage>101135</fpage><pub-id pub-id-type="doi">10.1016/j.lanepe.2024.101135</pub-id><pub-id pub-id-type="medline">39811787</pub-id></nlm-citation></ref><ref id="ref5"><label>5</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Goi</surname><given-names>PD</given-names> </name><name name-style="western"><surname>B&#x00FC;cker</surname><given-names>J</given-names> </name><name name-style="western"><surname>Vianna-Sulzbach</surname><given-names>M</given-names> </name><etal/></person-group><article-title>Pharmacological treatment and staging in bipolar disorder: evidence from clinical practice</article-title><source>Braz J Psychiatry</source><year>2015</year><volume>37</volume><issue>2</issue><fpage>121</fpage><lpage>125</lpage><pub-id pub-id-type="doi">10.1590/1516-4446-2014-1554</pub-id><pub-id pub-id-type="medline">26018648</pub-id></nlm-citation></ref><ref id="ref6"><label>6</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Crapanzano</surname><given-names>C</given-names> </name><name name-style="western"><surname>Casolaro</surname><given-names>I</given-names> </name><name name-style="western"><surname>Amendola</surname><given-names>C</given-names> </name><name name-style="western"><surname>Damiani</surname><given-names>S</given-names> </name></person-group><article-title>Lithium and valproate in bipolar disorder: from international evidence-based guidelines to clinical predictors</article-title><source>Clin Psychopharmacol Neurosci</source><year>2022</year><month>08</month><day>31</day><volume>20</volume><issue>3</issue><fpage>403</fpage><lpage>414</lpage><pub-id pub-id-type="doi">10.9758/cpn.2022.20.3.403</pub-id><pub-id pub-id-type="medline">35879025</pub-id></nlm-citation></ref><ref id="ref7"><label>7</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Bowden</surname><given-names>CL</given-names> </name></person-group><article-title>Anticonvulsants in bipolar disorders: current research and practice and future directions</article-title><source>Bipolar Disord</source><year>2009</year><month>06</month><volume>11 Suppl 2</volume><fpage>20</fpage><lpage>33</lpage><pub-id pub-id-type="doi">10.1111/j.1399-5618.2009.00708.x</pub-id><pub-id pub-id-type="medline">19538683</pub-id></nlm-citation></ref><ref id="ref8"><label>8</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Yalin</surname><given-names>N</given-names> </name><name name-style="western"><surname>Young</surname><given-names>AH</given-names> </name></person-group><article-title>Pharmacological treatment of bipolar depression: what are the current and emerging options?</article-title><source>Neuropsychiatr Dis Treat</source><year>2020</year><volume>16</volume><fpage>1459</fpage><lpage>1472</lpage><pub-id pub-id-type="doi">10.2147/NDT.S245166</pub-id><pub-id pub-id-type="medline">32606699</pub-id></nlm-citation></ref><ref id="ref9"><label>9</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Kowalczyk</surname><given-names>E</given-names> </name><name name-style="western"><surname>Koziej</surname><given-names>S</given-names> </name><name name-style="western"><surname>Soroka</surname><given-names>E</given-names> </name></person-group><article-title>Advances in mood disorder pharmacotherapy: evaluating new antipsychotics and mood stabilizers for bipolar disorder and schizophrenia</article-title><source>Med Sci Monit</source><year>2024</year><month>09</month><day>7</day><volume>30</volume><fpage>e945412</fpage><pub-id pub-id-type="doi">10.12659/MSM.945412</pub-id><pub-id pub-id-type="medline">39243127</pub-id></nlm-citation></ref><ref id="ref10"><label>10</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Patel</surname><given-names>R</given-names> </name><name name-style="western"><surname>Shetty</surname><given-names>H</given-names> </name><name name-style="western"><surname>Jackson</surname><given-names>R</given-names> </name><etal/></person-group><article-title>Delays before diagnosis and initiation of treatment in patients presenting to mental health services with bipolar disorder</article-title><source>PLoS One</source><year>2015</year><volume>10</volume><issue>5</issue><fpage>e0126530</fpage><pub-id pub-id-type="doi">10.1371/journal.pone.0126530</pub-id><pub-id pub-id-type="medline">25992560</pub-id></nlm-citation></ref><ref id="ref11"><label>11</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Davis</surname><given-names>J</given-names> </name><name name-style="western"><surname>Desmond</surname><given-names>M</given-names> </name><name name-style="western"><surname>Berk</surname><given-names>M</given-names> </name></person-group><article-title>Lithium and nephrotoxicity: a literature review of approaches to clinical management and risk stratification</article-title><source>BMC Nephrol</source><year>2018</year><month>11</month><day>3</day><volume>19</volume><issue>1</issue><fpage>305</fpage><pub-id pub-id-type="doi">10.1186/s12882-018-1101-4</pub-id><pub-id pub-id-type="medline">30390660</pub-id></nlm-citation></ref><ref id="ref12"><label>12</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Scigliano</surname><given-names>G</given-names> </name><name name-style="western"><surname>Ronchetti</surname><given-names>G</given-names> </name></person-group><article-title>Antipsychotic-induced metabolic and cardiovascular side effects in schizophrenia: a novel mechanistic hypothesis</article-title><source>CNS Drugs</source><year>2013</year><month>04</month><volume>27</volume><issue>4</issue><fpage>249</fpage><lpage>257</lpage><pub-id pub-id-type="doi">10.1007/s40263-013-0054-1</pub-id><pub-id pub-id-type="medline">23533011</pub-id></nlm-citation></ref><ref id="ref13"><label>13</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Goldberg</surname><given-names>JF</given-names> </name></person-group><article-title>Complex combination pharmacotherapy for bipolar disorder: knowing when less is more or more is better</article-title><source>FOC</source><year>2019</year><month>07</month><volume>17</volume><issue>3</issue><fpage>218</fpage><lpage>231</lpage><pub-id pub-id-type="doi">10.1176/appi.focus.20190008</pub-id></nlm-citation></ref><ref id="ref14"><label>14</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Bzdok</surname><given-names>D</given-names> </name><name name-style="western"><surname>Meyer-Lindenberg</surname><given-names>A</given-names> </name></person-group><article-title>Machine learning for precision psychiatry: opportunities and challenges</article-title><source>Biol Psychiatry Cogn Neurosci Neuroimaging</source><year>2018</year><month>03</month><volume>3</volume><issue>3</issue><fpage>223</fpage><lpage>230</lpage><pub-id pub-id-type="doi">10.1016/j.bpsc.2017.11.007</pub-id><pub-id pub-id-type="medline">29486863</pub-id></nlm-citation></ref><ref id="ref15"><label>15</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Chekroud</surname><given-names>AM</given-names> </name><name name-style="western"><surname>Zotti</surname><given-names>RJ</given-names> </name><name name-style="western"><surname>Shehzad</surname><given-names>Z</given-names> </name><etal/></person-group><article-title>Cross-trial prediction of treatment outcome in depression: a machine learning approach</article-title><source>Lancet Psychiatry</source><year>2016</year><month>03</month><volume>3</volume><issue>3</issue><fpage>243</fpage><lpage>250</lpage><pub-id pub-id-type="doi">10.1016/S2215-0366(15)00471-X</pub-id><pub-id pub-id-type="medline">26803397</pub-id></nlm-citation></ref><ref id="ref16"><label>16</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Passos</surname><given-names>IC</given-names> </name><name name-style="western"><surname>Mwangi</surname><given-names>B</given-names> </name><name name-style="western"><surname>Kapczinski</surname><given-names>F</given-names> </name></person-group><article-title>Big data analytics and machine learning: 2015 and beyond</article-title><source>Lancet Psychiatry</source><year>2016</year><month>01</month><volume>3</volume><issue>1</issue><fpage>13</fpage><lpage>15</lpage><pub-id pub-id-type="doi">10.1016/S2215-0366(15)00549-0</pub-id><pub-id pub-id-type="medline">26772057</pub-id></nlm-citation></ref><ref id="ref17"><label>17</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Iniesta</surname><given-names>R</given-names> </name><name name-style="western"><surname>Stahl</surname><given-names>D</given-names> </name><name name-style="western"><surname>McGuffin</surname><given-names>P</given-names> </name></person-group><article-title>Machine learning, statistical learning and the future of biological research in psychiatry</article-title><source>Psychol Med</source><year>2016</year><month>09</month><volume>46</volume><issue>12</issue><fpage>2455</fpage><lpage>2465</lpage><pub-id pub-id-type="doi">10.1017/S0033291716001367</pub-id><pub-id pub-id-type="medline">27406289</pub-id></nlm-citation></ref><ref id="ref18"><label>18</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Rotenberg</surname><given-names>L de S</given-names> </name><name name-style="western"><surname>Borges-J&#x00FA;nior</surname><given-names>RG</given-names> </name><name name-style="western"><surname>Lafer</surname><given-names>B</given-names> </name><name name-style="western"><surname>Salvini</surname><given-names>R</given-names> </name><name name-style="western"><surname>Dias</surname><given-names>R da S</given-names> </name></person-group><article-title>Exploring machine learning to predict depressive relapses of bipolar disorder patients</article-title><source>J Affect Disord</source><year>2021</year><month>12</month><day>1</day><volume>295</volume><fpage>681</fpage><lpage>687</lpage><pub-id pub-id-type="doi">10.1016/j.jad.2021.08.127</pub-id><pub-id pub-id-type="medline">34509784</pub-id></nlm-citation></ref><ref id="ref19"><label>19</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Cao</surname><given-names>P</given-names> </name><name name-style="western"><surname>Li</surname><given-names>R</given-names> </name><name name-style="western"><surname>Li</surname><given-names>Y</given-names> </name><etal/></person-group><article-title>Machine learning based differential diagnosis of schizophrenia, major depression disorder and bipolar disorder using structural magnetic resonance imaging</article-title><source>J Affect Disord</source><year>2025</year><month>08</month><day>15</day><volume>383</volume><fpage>20</fpage><lpage>31</lpage><pub-id pub-id-type="doi">10.1016/j.jad.2025.04.135</pub-id><pub-id pub-id-type="medline">40286928</pub-id></nlm-citation></ref><ref id="ref20"><label>20</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Torous</surname><given-names>J</given-names> </name><name name-style="western"><surname>Bucci</surname><given-names>S</given-names> </name><name name-style="western"><surname>Bell</surname><given-names>IH</given-names> </name><etal/></person-group><article-title>The growing field of digital psychiatry: current evidence and the future of apps, social media, chatbots, and virtual reality</article-title><source>World Psychiatry</source><year>2021</year><month>10</month><volume>20</volume><issue>3</issue><fpage>318</fpage><lpage>335</lpage><pub-id pub-id-type="doi">10.1002/wps.20883</pub-id><pub-id pub-id-type="medline">34505369</pub-id></nlm-citation></ref><ref id="ref21"><label>21</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Low</surname><given-names>DM</given-names> </name><name name-style="western"><surname>Bentley</surname><given-names>KH</given-names> </name><name name-style="western"><surname>Ghosh</surname><given-names>SS</given-names> </name></person-group><article-title>Automated assessment of psychiatric disorders using speech: a systematic review</article-title><source>Laryngoscope Investig Otolaryngol</source><year>2020</year><month>02</month><volume>5</volume><issue>1</issue><fpage>96</fpage><lpage>116</lpage><pub-id pub-id-type="doi">10.1002/lio2.354</pub-id><pub-id pub-id-type="medline">32128436</pub-id></nlm-citation></ref><ref id="ref22"><label>22</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Wu</surname><given-names>CT</given-names> </name><name name-style="western"><surname>Hsieh</surname><given-names>MH</given-names> </name><name name-style="western"><surname>Chen</surname><given-names>IM</given-names> </name><etal/></person-group><article-title>Using wearable device and machine learning to predict mood symptoms in bipolar disorder: development and usability study</article-title><source>JMIR Med Inform</source><year>2025</year><month>09</month><day>16</day><volume>13</volume><fpage>e66277</fpage><pub-id pub-id-type="doi">10.2196/66277</pub-id><pub-id pub-id-type="medline">40957006</pub-id></nlm-citation></ref><ref id="ref23"><label>23</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Beam</surname><given-names>AL</given-names> </name><name name-style="western"><surname>Kohane</surname><given-names>IS</given-names> </name></person-group><article-title>Big data and machine learning in health care</article-title><source>JAMA</source><year>2018</year><month>04</month><day>3</day><volume>319</volume><issue>13</issue><fpage>1317</fpage><lpage>1318</lpage><pub-id pub-id-type="doi">10.1001/jama.2017.18391</pub-id><pub-id pub-id-type="medline">29532063</pub-id></nlm-citation></ref><ref id="ref24"><label>24</label><nlm-citation citation-type="book"><person-group person-group-type="author"><name name-style="western"><surname>Gerke</surname><given-names>S</given-names> </name><name name-style="western"><surname>Minssen</surname><given-names>T</given-names> </name><name name-style="western"><surname>Cohen</surname><given-names>G</given-names> </name></person-group><article-title>Ethical and legal challenges of artificial intelligence-driven healthcare</article-title><source>Artificial Intelligence in Healthcare</source><year>2020</year><fpage>295</fpage><lpage>336</lpage><pub-id pub-id-type="doi">10.1016/B978-0-12-818438-7.00012-5</pub-id></nlm-citation></ref><ref id="ref25"><label>25</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Vayena</surname><given-names>E</given-names> </name><name name-style="western"><surname>Blasimme</surname><given-names>A</given-names> </name><name name-style="western"><surname>Cohen</surname><given-names>IG</given-names> </name></person-group><article-title>Machine learning in medicine: addressing ethical challenges</article-title><source>PLoS Med</source><year>2018</year><month>11</month><volume>15</volume><issue>11</issue><fpage>e1002689</fpage><pub-id pub-id-type="doi">10.1371/journal.pmed.1002689</pub-id><pub-id pub-id-type="medline">30399149</pub-id></nlm-citation></ref><ref id="ref26"><label>26</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Floridi</surname><given-names>L</given-names> </name><name name-style="western"><surname>Cowls</surname><given-names>J</given-names> </name><name name-style="western"><surname>Beltrametti</surname><given-names>M</given-names> </name><etal/></person-group><article-title>AI4People-an ethical framework for a good AI society: opportunities, risks, principles, and recommendations</article-title><source>Minds Mach (Dordr)</source><year>2018</year><volume>28</volume><issue>4</issue><fpage>689</fpage><lpage>707</lpage><pub-id pub-id-type="doi">10.1007/s11023-018-9482-5</pub-id><pub-id pub-id-type="medline">30930541</pub-id></nlm-citation></ref><ref id="ref27"><label>27</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Moons</surname><given-names>KGM</given-names> </name><name name-style="western"><surname>Damen</surname><given-names>JAA</given-names> </name><name name-style="western"><surname>Kaul</surname><given-names>T</given-names> </name><etal/></person-group><article-title>PROBAST+AI: an updated quality, risk of bias, and applicability assessment tool for prediction models using regression or artificial intelligence methods</article-title><source>BMJ</source><year>2025</year><month>03</month><day>24</day><volume>388</volume><fpage>e082505</fpage><pub-id pub-id-type="doi">10.1136/bmj-2024-082505</pub-id><pub-id pub-id-type="medline">40127903</pub-id></nlm-citation></ref><ref id="ref28"><label>28</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Ouzzani</surname><given-names>M</given-names> </name><name name-style="western"><surname>Hammady</surname><given-names>H</given-names> </name><name name-style="western"><surname>Fedorowicz</surname><given-names>Z</given-names> </name><name name-style="western"><surname>Elmagarmid</surname><given-names>A</given-names> </name></person-group><article-title>Rayyan-a web and mobile app for systematic reviews</article-title><source>Syst Rev</source><year>2016</year><month>12</month><day>5</day><volume>5</volume><issue>1</issue><fpage>210</fpage><pub-id pub-id-type="doi">10.1186/s13643-016-0384-4</pub-id><pub-id pub-id-type="medline">27919275</pub-id></nlm-citation></ref><ref id="ref29"><label>29</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>DerSimonian</surname><given-names>R</given-names> </name><name name-style="western"><surname>Laird</surname><given-names>N</given-names> </name></person-group><article-title>Meta-analysis in clinical trials</article-title><source>Control Clin Trials</source><year>1986</year><month>09</month><volume>7</volume><issue>3</issue><fpage>177</fpage><lpage>188</lpage><pub-id pub-id-type="doi">10.1016/0197-2456(86)90046-2</pub-id><pub-id pub-id-type="medline">3802833</pub-id></nlm-citation></ref><ref id="ref30"><label>30</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Agniel</surname><given-names>D</given-names> </name><name name-style="western"><surname>Normand</surname><given-names>SLT</given-names> </name><name name-style="western"><surname>Newcomer</surname><given-names>JW</given-names> </name><etal/></person-group><article-title>Revisiting diabetes risk of olanzapine versus aripiprazole in serious mental illness care</article-title><source>BJPsych Open</source><year>2024</year><month>08</month><day>8</day><volume>10</volume><issue>5</issue><fpage>e144</fpage><pub-id pub-id-type="doi">10.1192/bjo.2024.727</pub-id><pub-id pub-id-type="medline">39113461</pub-id></nlm-citation></ref><ref id="ref31"><label>31</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Brodeur</surname><given-names>S</given-names> </name><name name-style="western"><surname>Terrisse</surname><given-names>H</given-names> </name><name name-style="western"><surname>Pouchon</surname><given-names>A</given-names> </name><etal/></person-group><article-title>Pharmacological treatment profiles in the FACE-BD cohort: an unsupervised machine learning study, applied to a nationwide bipolar cohort<sup>&#x2730;</sup></article-title><source>J Affect Disord</source><year>2021</year><month>05</month><day>1</day><volume>286</volume><fpage>309</fpage><lpage>319</lpage><pub-id pub-id-type="doi">10.1016/j.jad.2021.02.036</pub-id><pub-id pub-id-type="medline">33770539</pub-id></nlm-citation></ref><ref id="ref32"><label>32</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Cearns</surname><given-names>M</given-names> </name><name name-style="western"><surname>Amare</surname><given-names>AT</given-names> </name><name name-style="western"><surname>Schubert</surname><given-names>KO</given-names> </name><etal/></person-group><article-title>Using polygenic scores and clinical data for bipolar disorder patient stratification and lithium response prediction: machine learning approach</article-title><source>Br J Psychiatry</source><year>2022</year><month>04</month><volume>220</volume><issue>4</issue><fpage>219</fpage><lpage>228</lpage><pub-id pub-id-type="doi">10.1192/bjp.2022.28</pub-id><pub-id pub-id-type="medline">35225756</pub-id></nlm-citation></ref><ref id="ref33"><label>33</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>D&#x00ED;az-Zuluaga</surname><given-names>AM</given-names> </name><name name-style="western"><surname>V&#x00E9;lez</surname><given-names>JI</given-names> </name><name name-style="western"><surname>Cuartas</surname><given-names>M</given-names> </name><etal/></person-group><article-title>Ancestry component as a major predictor of lithium response in the treatment of bipolar disorder</article-title><source>J Affect Disord</source><year>2023</year><month>07</month><day>1</day><volume>332</volume><fpage>203</fpage><lpage>209</lpage><pub-id pub-id-type="doi">10.1016/j.jad.2023.03.058</pub-id><pub-id pub-id-type="medline">36997125</pub-id></nlm-citation></ref><ref id="ref34"><label>34</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Edgcomb</surname><given-names>JB</given-names> </name><name name-style="western"><surname>Shaddox</surname><given-names>T</given-names> </name><name name-style="western"><surname>Hellemann</surname><given-names>G</given-names> </name><name name-style="western"><surname>Brooks</surname><given-names>JO</given-names>  <suffix>III</suffix></name></person-group><article-title>Predicting suicidal behavior and self-harm after general hospitalization of adults with serious mental illness</article-title><source>J Psychiatr Res</source><year>2021</year><month>04</month><volume>136</volume><fpage>515</fpage><lpage>521</lpage><pub-id pub-id-type="doi">10.1016/j.jpsychires.2020.10.024</pub-id><pub-id pub-id-type="medline">33218748</pub-id></nlm-citation></ref><ref id="ref35"><label>35</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Fleck</surname><given-names>DE</given-names> </name><name name-style="western"><surname>Ernest</surname><given-names>N</given-names> </name><name name-style="western"><surname>Adler</surname><given-names>CM</given-names> </name><etal/></person-group><article-title>Prediction of lithium response in first-episode mania using the LITHium Intelligent Agent (LITHIA): Pilot data and proof-of-concept</article-title><source>Bipolar Disord</source><year>2017</year><month>06</month><volume>19</volume><issue>4</issue><fpage>259</fpage><lpage>272</lpage><pub-id pub-id-type="doi">10.1111/bdi.12507</pub-id><pub-id pub-id-type="medline">28574156</pub-id></nlm-citation></ref><ref id="ref36"><label>36</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Hayes</surname><given-names>JF</given-names> </name><name name-style="western"><surname>Ben Abdesslem</surname><given-names>F</given-names> </name><name name-style="western"><surname>Eloranta</surname><given-names>S</given-names> </name><name name-style="western"><surname>Osborn</surname><given-names>DPJ</given-names> </name><name name-style="western"><surname>Boman</surname><given-names>M</given-names> </name></person-group><article-title>Predicting maintenance lithium response for bipolar disorder from electronic health records-a retrospective study</article-title><source>PeerJ</source><year>2024</year><volume>12</volume><fpage>e17841</fpage><pub-id pub-id-type="doi">10.7717/peerj.17841</pub-id><pub-id pub-id-type="medline">39421428</pub-id></nlm-citation></ref><ref id="ref37"><label>37</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Kim</surname><given-names>TT</given-names> </name><name name-style="western"><surname>Dufour</surname><given-names>S</given-names> </name><name name-style="western"><surname>Xu</surname><given-names>C</given-names> </name><etal/></person-group><article-title>Predictive modeling for response to lithium and quetiapine in bipolar disorder</article-title><source>Bipolar Disord</source><year>2019</year><month>08</month><volume>21</volume><issue>5</issue><fpage>428</fpage><lpage>436</lpage><pub-id pub-id-type="doi">10.1111/bdi.12752</pub-id><pub-id pub-id-type="medline">30729637</pub-id></nlm-citation></ref><ref id="ref38"><label>38</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Lei</surname><given-names>D</given-names> </name><name name-style="western"><surname>Li</surname><given-names>W</given-names> </name><name name-style="western"><surname>Tallman</surname><given-names>MJ</given-names> </name><etal/></person-group><article-title>Changes in the structural brain connectome over the course of a nonrandomized clinical trial for acute mania</article-title><source>Neuropsychopharmacol</source><year>2022</year><month>10</month><volume>47</volume><issue>11</issue><fpage>1961</fpage><lpage>1968</lpage><pub-id pub-id-type="doi">10.1038/s41386-022-01328-y</pub-id></nlm-citation></ref><ref id="ref39"><label>39</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Lieslehto</surname><given-names>J</given-names> </name><name name-style="western"><surname>Tiihonen</surname><given-names>J</given-names> </name><name name-style="western"><surname>L&#x00E4;hteenvuo</surname><given-names>M</given-names> </name><etal/></person-group><article-title>Machine learning-based mortality risk assessment in first-episode bipolar disorder: a transdiagnostic external validation study</article-title><source>EClinicalMedicine</source><year>2025</year><month>03</month><volume>81</volume><fpage>103108</fpage><pub-id pub-id-type="doi">10.1016/j.eclinm.2025.103108</pub-id><pub-id pub-id-type="medline">40034574</pub-id></nlm-citation></ref><ref id="ref40"><label>40</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Lieslehto</surname><given-names>J</given-names> </name><name name-style="western"><surname>Tiihonen</surname><given-names>J</given-names> </name><name name-style="western"><surname>L&#x00E4;hteenvuo</surname><given-names>M</given-names> </name><etal/></person-group><article-title>Relapse risk prediction in patients with first-episode bipolar disorder: development, external validation, and pharmacotherapy associations of a machine learning model</article-title><source>Mol Psychiatry</source><year>2025</year><month>12</month><volume>30</volume><issue>12</issue><fpage>5722</fpage><lpage>5730</lpage><pub-id pub-id-type="doi">10.1038/s41380-025-03316-2</pub-id><pub-id pub-id-type="medline">41131281</pub-id></nlm-citation></ref><ref id="ref41"><label>41</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Lv</surname><given-names>D</given-names> </name><name name-style="western"><surname>Yan</surname><given-names>HH</given-names> </name><name name-style="western"><surname>Zhang</surname><given-names>CG</given-names> </name><etal/></person-group><article-title>Kcc-ReHo and Cohe-ReHo in bipolar disorder: their associated genes and potential for diagnosis and treatment prediction</article-title><source>Neuropharmacology</source><year>2025</year><month>11</month><day>1</day><volume>278</volume><fpage>110575</fpage><pub-id pub-id-type="doi">10.1016/j.neuropharm.2025.110575</pub-id><pub-id pub-id-type="medline">40578678</pub-id></nlm-citation></ref><ref id="ref42"><label>42</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Marie-Claire</surname><given-names>C</given-names> </name><name name-style="western"><surname>Courtin</surname><given-names>C</given-names> </name><name name-style="western"><surname>Bellivier</surname><given-names>F</given-names> </name><name name-style="western"><surname>Scott</surname><given-names>J</given-names> </name><name name-style="western"><surname>Etain</surname><given-names>B</given-names> </name></person-group><article-title>Methylomic biomarkers of lithium response in bipolar disorder: a proof of transferability study</article-title><source>Pharmaceuticals (Basel)</source><year>2022</year><month>01</month><day>23</day><volume>15</volume><issue>2</issue><fpage>133</fpage><pub-id pub-id-type="doi">10.3390/ph15020133</pub-id><pub-id pub-id-type="medline">35215246</pub-id></nlm-citation></ref><ref id="ref43"><label>43</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Mizrahi</surname><given-names>L</given-names> </name><name name-style="western"><surname>Choudhary</surname><given-names>A</given-names> </name><name name-style="western"><surname>Ofer</surname><given-names>P</given-names> </name><etal/></person-group><article-title>Immunoglobulin genes expressed in lymphoblastoid cell lines discern and predict lithium response in bipolar disorder patients</article-title><source>Mol Psychiatry</source><year>2023</year><month>10</month><volume>28</volume><issue>10</issue><fpage>4280</fpage><lpage>4293</lpage><pub-id pub-id-type="doi">10.1038/s41380-023-02183-z</pub-id><pub-id pub-id-type="medline">37488168</pub-id></nlm-citation></ref><ref id="ref44"><label>44</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Mora</surname><given-names>F</given-names> </name><name name-style="western"><surname>G&#x00F3;mez S&#x00E1;nchez-Lafuente</surname><given-names>C</given-names> </name><name name-style="western"><surname>De Iceta</surname><given-names>M</given-names> </name><etal/></person-group><article-title>Lurasidone uses and dosages in Spain: RETROLUR, a real-world retrospective analysis using artificial intelligence</article-title><source>Front Psychiatry</source><year>2024</year><volume>15</volume><fpage>1506142</fpage><pub-id pub-id-type="doi">10.3389/fpsyt.2024.1506142</pub-id><pub-id pub-id-type="medline">40013022</pub-id></nlm-citation></ref><ref id="ref45"><label>45</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Nestsiarovich</surname><given-names>A</given-names> </name><name name-style="western"><surname>Kumar</surname><given-names>P</given-names> </name><name name-style="western"><surname>Lauve</surname><given-names>NR</given-names> </name><etal/></person-group><article-title>Using machine learning imputed outcomes to assess drug-dependent risk of self-harm in patients with bipolar disorder: a comparative effectiveness study</article-title><source>JMIR Ment Health</source><year>2021</year><month>04</month><day>21</day><volume>8</volume><issue>4</issue><fpage>e24522</fpage><pub-id pub-id-type="doi">10.2196/24522</pub-id><pub-id pub-id-type="medline">33688834</pub-id></nlm-citation></ref><ref id="ref46"><label>46</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Nielsen</surname><given-names>SFV</given-names> </name><name name-style="western"><surname>Madsen</surname><given-names>KH</given-names> </name><name name-style="western"><surname>Vinberg</surname><given-names>M</given-names> </name><name name-style="western"><surname>Kessing</surname><given-names>LV</given-names> </name><name name-style="western"><surname>Siebner</surname><given-names>HR</given-names> </name><name name-style="western"><surname>Miskowiak</surname><given-names>KW</given-names> </name></person-group><article-title>Whole-brain exploratory analysis of functional task response following erythropoietin treatment in mood disorders: a supervised machine learning approach</article-title><source>Front Neurosci</source><year>2019</year><volume>13</volume><fpage>1246</fpage><pub-id pub-id-type="doi">10.3389/fnins.2019.01246</pub-id><pub-id pub-id-type="medline">31824247</pub-id></nlm-citation></ref><ref id="ref47"><label>47</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Nunes</surname><given-names>A</given-names> </name><name name-style="western"><surname>Ardau</surname><given-names>R</given-names> </name><name name-style="western"><surname>Bergh&#x00F6;fer</surname><given-names>A</given-names> </name><etal/></person-group><article-title>Prediction of lithium response using clinical data</article-title><source>Acta Psychiatr Scand</source><year>2020</year><month>02</month><volume>141</volume><issue>2</issue><fpage>131</fpage><lpage>141</lpage><pub-id pub-id-type="doi">10.1111/acps.13122</pub-id><pub-id pub-id-type="medline">31667829</pub-id></nlm-citation></ref><ref id="ref48"><label>48</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Palau</surname><given-names>P</given-names> </name><name name-style="western"><surname>Solanes</surname><given-names>A</given-names> </name><name name-style="western"><surname>Madre</surname><given-names>M</given-names> </name><etal/></person-group><article-title>Improved estimation of the risk of manic relapse by combining clinical and brain scan data</article-title><source>Span J Psychiatry Ment Health</source><year>2023</year><volume>16</volume><issue>4</issue><fpage>235</fpage><lpage>243</lpage><pub-id pub-id-type="doi">10.1016/j.rpsm.2023.01.001</pub-id><pub-id pub-id-type="medline">37839962</pub-id></nlm-citation></ref><ref id="ref49"><label>49</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Pettorruso</surname><given-names>M</given-names> </name><name name-style="western"><surname>Guidotti</surname><given-names>R</given-names> </name><name name-style="western"><surname>d&#x2019;Andrea</surname><given-names>G</given-names> </name><etal/></person-group><article-title>Predicting outcome with Intranasal Esketamine treatment: a machine-LEARNING, three-month study in treatment-resistant depression (ESK-LEARNING)</article-title><source>Psychiatry Res</source><year>2023</year><month>09</month><volume>327</volume><fpage>115378</fpage><pub-id pub-id-type="doi">10.1016/j.psychres.2023.115378</pub-id><pub-id pub-id-type="medline">37574600</pub-id></nlm-citation></ref><ref id="ref50"><label>50</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Poulos</surname><given-names>J</given-names> </name><name name-style="western"><surname>Horvitz-Lennon</surname><given-names>M</given-names> </name><name name-style="western"><surname>Zelevinsky</surname><given-names>K</given-names> </name><etal/></person-group><article-title>Targeted learning in observational studies with multi-valued treatments: an evaluation of antipsychotic drug treatment safety</article-title><source>Stat Med</source><year>2024</year><month>04</month><day>15</day><volume>43</volume><issue>8</issue><fpage>1489</fpage><lpage>1508</lpage><pub-id pub-id-type="doi">10.1002/sim.10003</pub-id><pub-id pub-id-type="medline">38314950</pub-id></nlm-citation></ref><ref id="ref51"><label>51</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Ross</surname><given-names>EL</given-names> </name><name name-style="western"><surname>Bossarte</surname><given-names>RM</given-names> </name><name name-style="western"><surname>Dobscha</surname><given-names>SK</given-names> </name><etal/></person-group><article-title>Estimated average treatment effect of psychiatric hospitalization in patients with suicidal behaviors: a precision treatment analysis</article-title><source>JAMA Psychiatry</source><year>2024</year><month>02</month><day>1</day><volume>81</volume><issue>2</issue><fpage>135</fpage><lpage>143</lpage><pub-id pub-id-type="doi">10.1001/jamapsychiatry.2023.3994</pub-id><pub-id pub-id-type="medline">37851457</pub-id></nlm-citation></ref><ref id="ref52"><label>52</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Ryan</surname><given-names>MC</given-names> </name><name name-style="western"><surname>Hong</surname><given-names>LE</given-names> </name><name name-style="western"><surname>Hatch</surname><given-names>KS</given-names> </name><etal/></person-group><article-title>The additive impact of cardio-metabolic disorders and psychiatric illnesses on accelerated brain aging</article-title><source>Hum Brain Mapp</source><year>2022</year><month>04</month><day>15</day><volume>43</volume><issue>6</issue><fpage>1997</fpage><lpage>2010</lpage><pub-id pub-id-type="doi">10.1002/hbm.25769</pub-id><pub-id pub-id-type="medline">35112422</pub-id></nlm-citation></ref><ref id="ref53"><label>53</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Salari</surname><given-names>N</given-names> </name><name name-style="western"><surname>Pilangorgi</surname><given-names>SS</given-names> </name><name name-style="western"><surname>Almasi</surname><given-names>A</given-names> </name><name name-style="western"><surname>Shahsavari</surname><given-names>S</given-names> </name><name name-style="western"><surname>Fournier</surname><given-names>AJ</given-names> </name></person-group><article-title>Classification of patients with lithium-treated bipolar disorder based on gene expression: Dirichlet Bayesian network model</article-title><source>Egypt J Med Hum Genet</source><year>2025</year><volume>26</volume><issue>1</issue><fpage>64</fpage><pub-id pub-id-type="doi">10.1186/s43042-025-00690-y</pub-id></nlm-citation></ref><ref id="ref54"><label>54</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Salvini</surname><given-names>R</given-names> </name><name name-style="western"><surname>da Silva Dias</surname><given-names>R</given-names> </name><name name-style="western"><surname>Lafer</surname><given-names>B</given-names> </name><name name-style="western"><surname>Dutra</surname><given-names>I</given-names> </name></person-group><article-title>A multi-relational model for depression relapse in patients with bipolar disorder</article-title><source>Stud Health Technol Inform</source><year>2015</year><volume>216</volume><issue>741-5</issue><fpage>741</fpage><lpage>745</lpage><pub-id pub-id-type="medline">26262150</pub-id></nlm-citation></ref><ref id="ref55"><label>55</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Scott</surname><given-names>J</given-names> </name><name name-style="western"><surname>Lajnef</surname><given-names>M</given-names> </name><name name-style="western"><surname>Icick</surname><given-names>R</given-names> </name><name name-style="western"><surname>Bellivier</surname><given-names>F</given-names> </name><name name-style="western"><surname>Marie-Claire</surname><given-names>C</given-names> </name><name name-style="western"><surname>Etain</surname><given-names>B</given-names> </name></person-group><article-title>A comparison of different approaches to clinical phenotyping of lithium response: a proof of principle study employing genetic variants of three candidate circadian genes</article-title><source>Pharmaceuticals (Basel)</source><year>2021</year><month>10</month><day>23</day><volume>14</volume><issue>11</issue><fpage>1072</fpage><pub-id pub-id-type="doi">10.3390/ph14111072</pub-id><pub-id pub-id-type="medline">34832854</pub-id></nlm-citation></ref><ref id="ref56"><label>56</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Stone</surname><given-names>W</given-names> </name><name name-style="western"><surname>Nunes</surname><given-names>A</given-names> </name><name name-style="western"><surname>Akiyama</surname><given-names>K</given-names> </name><etal/></person-group><article-title>Prediction of lithium response using genomic data</article-title><source>Sci Rep</source><year>2021</year><month>01</month><day>13</day><volume>11</volume><issue>1</issue><fpage>1155</fpage><pub-id pub-id-type="doi">10.1038/s41598-020-80814-z</pub-id><pub-id pub-id-type="medline">33441847</pub-id></nlm-citation></ref><ref id="ref57"><label>57</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Tripathi</surname><given-names>U</given-names> </name><name name-style="western"><surname>Mizrahi</surname><given-names>L</given-names> </name><name name-style="western"><surname>Alda</surname><given-names>M</given-names> </name><name name-style="western"><surname>Falkovich</surname><given-names>G</given-names> </name><name name-style="western"><surname>Stern</surname><given-names>S</given-names> </name></person-group><article-title>Information theory characteristics improve the prediction of lithium response in bipolar disorder patients using a support vector machine classifier</article-title><source>Bipolar Disord</source><year>2023</year><month>03</month><volume>25</volume><issue>2</issue><fpage>110</fpage><lpage>127</lpage><pub-id pub-id-type="doi">10.1111/bdi.13282</pub-id><pub-id pub-id-type="medline">36479788</pub-id></nlm-citation></ref><ref id="ref58"><label>58</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Van Gestel</surname><given-names>H</given-names> </name><name name-style="western"><surname>Franke</surname><given-names>K</given-names> </name><name name-style="western"><surname>Petite</surname><given-names>J</given-names> </name><etal/></person-group><article-title>Brain age in bipolar disorders: effects of lithium treatment</article-title><source>Aust N Z J Psychiatry</source><year>2019</year><month>12</month><volume>53</volume><issue>12</issue><fpage>1179</fpage><lpage>1188</lpage><pub-id pub-id-type="doi">10.1177/0004867419857814</pub-id><pub-id pub-id-type="medline">31244332</pub-id></nlm-citation></ref><ref id="ref59"><label>59</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Xi</surname><given-names>C</given-names> </name><name name-style="western"><surname>Lu</surname><given-names>B</given-names> </name><name name-style="western"><surname>Guo</surname><given-names>X</given-names> </name><name name-style="western"><surname>Qin</surname><given-names>Z</given-names> </name><name name-style="western"><surname>Yan</surname><given-names>C</given-names> </name><name name-style="western"><surname>Hu</surname><given-names>S</given-names> </name></person-group><article-title>Characteristics of brain network connectome and connectome-based efficacy predictive model in bipolar depression</article-title><source>Mol Psychiatry</source><year>2025</year><month>11</month><volume>30</volume><issue>11</issue><fpage>5150</fpage><lpage>5160</lpage><pub-id pub-id-type="doi">10.1038/s41380-025-03099-6</pub-id><pub-id pub-id-type="medline">40615558</pub-id></nlm-citation></ref><ref id="ref60"><label>60</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Zhang</surname><given-names>J</given-names> </name><name name-style="western"><surname>Wang</surname><given-names>Y</given-names> </name><name name-style="western"><surname>Zhang</surname><given-names>W</given-names> </name><etal/></person-group><article-title>The development and validation of a prediction model of lithium carbonate blood concentration by artificial neural network: a retrospective study</article-title><source>Ann Palliat Med</source><year>2022</year><month>12</month><volume>11</volume><issue>12</issue><fpage>3718</fpage><lpage>3726</lpage><pub-id pub-id-type="doi">10.21037/apm-22-1237</pub-id><pub-id pub-id-type="medline">36635997</pub-id></nlm-citation></ref><ref id="ref61"><label>61</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Zhang</surname><given-names>L</given-names> </name><name name-style="western"><surname>Yan</surname><given-names>H</given-names> </name><name name-style="western"><surname>Zhang</surname><given-names>C</given-names> </name><etal/></person-group><article-title>The application of amplitude of low-frequency fluctuations metrics in the diagnosis and prediction of treatment response as well as their associated genes and biological processes in patients with bipolar disorder</article-title><source>Transl Psychiatry</source><year>2025</year><month>10</month><day>31</day><volume>15</volume><issue>1</issue><fpage>446</fpage><pub-id pub-id-type="doi">10.1038/s41398-025-03673-0</pub-id><pub-id pub-id-type="medline">41173827</pub-id></nlm-citation></ref><ref id="ref62"><label>62</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Zheng</surname><given-names>P</given-names> </name><name name-style="western"><surname>Yu</surname><given-names>Z</given-names> </name><name name-style="western"><surname>Mo</surname><given-names>L</given-names> </name><etal/></person-group><article-title>An individualized medication model of sodium valproate for patients with bipolar disorder based on machine learning and deep learning techniques</article-title><source>Front Pharmacol</source><year>2022</year><volume>13</volume><fpage>890221</fpage><pub-id pub-id-type="doi">10.3389/fphar.2022.890221</pub-id><pub-id pub-id-type="medline">36339624</pub-id></nlm-citation></ref><ref id="ref63"><label>63</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Zhu</surname><given-names>X</given-names> </name><name name-style="western"><surname>Hu</surname><given-names>J</given-names> </name><name name-style="western"><surname>Xiao</surname><given-names>T</given-names> </name><name name-style="western"><surname>Huang</surname><given-names>S</given-names> </name><name name-style="western"><surname>Shang</surname><given-names>D</given-names> </name><name name-style="western"><surname>Wen</surname><given-names>Y</given-names> </name></person-group><article-title>Integrating machine learning with electronic health record data to facilitate detection of prolactin level and pharmacovigilance signals in olanzapine-treated patients</article-title><source>Front Endocrinol</source><year>2022</year><month>10</month><day>13</day><volume>13</volume><fpage>1011492</fpage><pub-id pub-id-type="doi">10.3389/fendo.2022.1011492</pub-id></nlm-citation></ref><ref id="ref64"><label>64</label><nlm-citation citation-type="web"><article-title>Fehmi 8/lithium</article-title><source>GitHub</source><access-date>2026-07-06</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://github.com/fehmi8/lithium">https://github.com/fehmi8/lithium</ext-link></comment></nlm-citation></ref><ref id="ref65"><label>65</label><nlm-citation citation-type="web"><article-title>MIRACLE-FEP; mortality risk assessment calculator for early first-episode psychosis</article-title><source>Shinyapps.io</source><access-date>2026-07-06</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://johannes-lieslehto.shinyapps.io/miracle-fep">https://johannes-lieslehto.shinyapps.io/miracle-fep</ext-link></comment></nlm-citation></ref><ref id="ref66"><label>66</label><nlm-citation citation-type="web"><person-group person-group-type="author"><name name-style="western"><surname>Lieslehto</surname><given-names>J</given-names> </name></person-group><article-title>FEBD_relapse_ML</article-title><source>GitHub</source><access-date>2026-07-06</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://github.com/johpulkk/FEBD_relapse_ML">https://github.com/johpulkk/FEBD_relapse_ML</ext-link></comment></nlm-citation></ref><ref id="ref67"><label>67</label><nlm-citation citation-type="web"><article-title>Lithum-respose-predictor [computer software</article-title><source>GitHub</source><access-date>2026-07-06</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://github.com/Precision-Disease-Modeling-Lab/Lithum-Respose-Predictor">https://github.com/Precision-Disease-Modeling-Lab/Lithum-Respose-Predictor</ext-link></comment></nlm-citation></ref><ref id="ref68"><label>68</label><nlm-citation citation-type="web"><article-title>Self harm pharmacotherapy psychotherapy</article-title><source>GitLab</source><access-date>2026-07-06</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://gitlab.com/PCORIUNMPUBLIC/self_harm_pharmacotherapy_psychotherapy">https://gitlab.com/PCORIUNMPUBLIC/self_harm_pharmacotherapy_psychotherapy</ext-link></comment></nlm-citation></ref><ref id="ref69"><label>69</label><nlm-citation citation-type="web"><article-title>MRIPredict mania</article-title><source>MRIPredict</source><access-date>2026-07-06</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://www.mripredict.com/mania/">https://www.mripredict.com/mania/</ext-link></comment></nlm-citation></ref><ref id="ref70"><label>70</label><nlm-citation citation-type="web"><article-title>Jvpoulos/multi-tmle</article-title><source>GitHub</source><access-date>2026-07-06</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://github.com/jvpoulos/multi-tmle">https://github.com/jvpoulos/multi-tmle</ext-link></comment></nlm-citation></ref><ref id="ref71"><label>71</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Brunn</surname><given-names>M</given-names> </name><name name-style="western"><surname>Diefenbacher</surname><given-names>A</given-names> </name><name name-style="western"><surname>Courtet</surname><given-names>P</given-names> </name><name name-style="western"><surname>Genieys</surname><given-names>W</given-names> </name></person-group><article-title>The future is knocking: how artificial intelligence will fundamentally change psychiatry</article-title><source>Acad Psychiatry</source><year>2020</year><month>08</month><volume>44</volume><issue>4</issue><fpage>461</fpage><lpage>466</lpage><pub-id pub-id-type="doi">10.1007/s40596-020-01243-8</pub-id><pub-id pub-id-type="medline">32424706</pub-id></nlm-citation></ref><ref id="ref72"><label>72</label><nlm-citation citation-type="web"><article-title>Regulation (EU) 2024/1689 of the european parliament and of the council of 13 june 2024 laying down harmonised rules on artificial intelligence and amending regulations (EC) no 300/2008, (EU) no 167/2013, (EU) no 168/2013, (EU) 2018/858, (EU) 2018/1139 and (EU) 2019/2144 and directives 2014/90/EU, (EU) 2016/797 and (EU) 2020/1828 (artificial intelligence act)</article-title><source>EUR-Lex</source><access-date>2026-07-10</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://eur-lex.europa.EU/legal-content/EN/TXT/?qid=1760706020411&#x0026;uri=CELEX%3A32024R1689">https://eur-lex.europa.EU/legal-content/EN/TXT/?qid=1760706020411&#x0026;uri=CELEX%3A32024R1689</ext-link></comment></nlm-citation></ref><ref id="ref73"><label>73</label><nlm-citation citation-type="web"><article-title>Reflection paper on the use of artificial intelligence (AI) in the medicinal product lifecycle</article-title><source>European medicines agency (EMA)</source><year>2023</year><access-date>2026-06-29</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://www.ema.europa.eu/en/documents/scientific-guideline/reflection-paper-use-artificial-intelligence-ai-medicinal-product-lifecycle_en.pdf">https://www.ema.europa.eu/en/documents/scientific-guideline/reflection-paper-use-artificial-intelligence-ai-medicinal-product-lifecycle_en.pdf</ext-link></comment></nlm-citation></ref><ref id="ref74"><label>74</label><nlm-citation citation-type="web"><article-title>Artificial intelligence&#x2011;enabled device software functions: lifecycle management and marketing submission recommendations</article-title><source>US Food and Drug Administration (FDA)</source><year>2025</year><access-date>2026-07-10</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://www.fda.gov/regulatory-information/search-fda-guidance-documents/artificial-intelligence-enabled-device-software-functions-lifecycle-management-and-marketing">https://www.fda.gov/regulatory-information/search-fda-guidance-documents/artificial-intelligence-enabled-device-software-functions-lifecycle-management-and-marketing</ext-link></comment></nlm-citation></ref><ref id="ref75"><label>75</label><nlm-citation citation-type="web"><article-title>Good machine learning practice for medical device development: guiding principles</article-title><source>US Food and Drug Administration (FDA)</source><access-date>2026-07-10</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://www.fda.gov/medical-devices/software-medical-device-samd/good-machine-learning-practice-medical-device-development-guiding-principles">https://www.fda.gov/medical-devices/software-medical-device-samd/good-machine-learning-practice-medical-device-development-guiding-principles</ext-link></comment></nlm-citation></ref><ref id="ref76"><label>76</label><nlm-citation citation-type="journal"><article-title>FUTURE-AI: international consensus guideline for trustworthy and deployable artificial intelligence in healthcare</article-title><source>BMJ</source><year>2025</year><month>02</month><day>17</day><volume>388</volume><fpage>r340</fpage><pub-id pub-id-type="doi">10.1136/bmj.r340</pub-id><pub-id pub-id-type="medline">39961614</pub-id></nlm-citation></ref><ref id="ref77"><label>77</label><nlm-citation citation-type="web"><article-title>Regulation (EU) 2017/745 of the european parliament and of the council of 5 april 2017 on medical devices, amending directive 2001/83/EC, regulation (EC) no 178/2002 and regulation (EC) no 1223/2009 and repealing council directives 90/385/EEC and 93/42/EEC (text with EEA relevance)</article-title><source>EUR-Lex</source><year>2017</year><access-date>2026-07-10</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://eur-lex.europa.eu/legal-content/EN/TXT/?uri=celex%253A32017R0745">https://eur-lex.europa.eu/legal-content/EN/TXT/?uri=celex%253A32017R0745</ext-link></comment></nlm-citation></ref><ref id="ref78"><label>78</label><nlm-citation citation-type="web"><article-title>Regulation (EU) 2017/746 of the european parliament and of the council of 5 april 2017 on in vitro diagnostic medical devices and repealing directive 98/79/EC and commission decision 2010/227/EU (text with EEA relevance)</article-title><source>EUR-Lex</source><access-date>2026-07-10</access-date><comment><ext-link ext-link-type="uri" xlink:href="https://eur-lex.europa.eu/legal-content/EN/TXT/?uri=celex%3A32017R0746">https://eur-lex.europa.eu/legal-content/EN/TXT/?uri=celex%3A32017R0746</ext-link></comment></nlm-citation></ref><ref id="ref79"><label>79</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Shatte</surname><given-names>ABR</given-names> </name><name name-style="western"><surname>Hutchinson</surname><given-names>DM</given-names> </name><name name-style="western"><surname>Teague</surname><given-names>SJ</given-names> </name></person-group><article-title>Machine learning in mental health: a scoping review of methods and applications</article-title><source>Psychol Med</source><year>2019</year><month>07</month><volume>49</volume><issue>9</issue><fpage>1426</fpage><lpage>1448</lpage><pub-id pub-id-type="doi">10.1017/S0033291719000151</pub-id><pub-id pub-id-type="medline">30744717</pub-id></nlm-citation></ref><ref id="ref80"><label>80</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Dwyer</surname><given-names>DB</given-names> </name><name name-style="western"><surname>Falkai</surname><given-names>P</given-names> </name><name name-style="western"><surname>Koutsouleris</surname><given-names>N</given-names> </name></person-group><article-title>Machine learning approaches for clinical psychology and psychiatry</article-title><source>Annu Rev Clin Psychol</source><year>2018</year><month>05</month><day>7</day><volume>14</volume><fpage>91</fpage><lpage>118</lpage><pub-id pub-id-type="doi">10.1146/annurev-clinpsy-032816-045037</pub-id><pub-id pub-id-type="medline">29401044</pub-id></nlm-citation></ref><ref id="ref81"><label>81</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Mohr</surname><given-names>DC</given-names> </name><name name-style="western"><surname>Zhang</surname><given-names>M</given-names> </name><name name-style="western"><surname>Schueller</surname><given-names>SM</given-names> </name></person-group><article-title>Personal sensing: understanding mental health using ubiquitous sensors and machine learning</article-title><source>Annu Rev Clin Psychol</source><year>2017</year><month>05</month><day>8</day><volume>13</volume><fpage>23</fpage><lpage>47</lpage><pub-id pub-id-type="doi">10.1146/annurev-clinpsy-032816-044949</pub-id><pub-id pub-id-type="medline">28375728</pub-id></nlm-citation></ref><ref id="ref82"><label>82</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Walsh</surname><given-names>CG</given-names> </name><name name-style="western"><surname>Ribeiro</surname><given-names>JD</given-names> </name><name name-style="western"><surname>Franklin</surname><given-names>JC</given-names> </name></person-group><article-title>Predicting risk of suicide attempts over time through machine learning</article-title><source>Clin Psychol Sci</source><year>2017</year><month>05</month><volume>5</volume><issue>3</issue><fpage>457</fpage><lpage>469</lpage><pub-id pub-id-type="doi">10.1177/2167702617691560</pub-id></nlm-citation></ref><ref id="ref83"><label>83</label><nlm-citation citation-type="journal"><person-group person-group-type="author"><name name-style="western"><surname>Belsher</surname><given-names>BE</given-names> </name><name name-style="western"><surname>Smolenski</surname><given-names>DJ</given-names> </name><name name-style="western"><surname>Pruitt</surname><given-names>LD</given-names> </name><etal/></person-group><article-title>Prediction models for suicide attempts and deaths: a systematic review and simulation</article-title><source>JAMA Psychiatry</source><year>2019</year><month>06</month><day>1</day><volume>76</volume><issue>6</issue><fpage>642</fpage><lpage>651</lpage><pub-id pub-id-type="doi">10.1001/jamapsychiatry.2019.0174</pub-id><pub-id pub-id-type="medline">30865249</pub-id></nlm-citation></ref></ref-list><app-group><supplementary-material id="app1"><label>Multimedia Appendix 1</label><p>Definitions of technical terms, the search strategies used across 4 electronic databases, the quality, risk of bias, and applicability assessments of the included studies conducted using PROBAST-AI, the objectives of the 35 included studies, and the dataset used to construct the bubble chart.</p><media xlink:href="mental_v13i1e93307_app1.docx" xlink:title="DOCX File, 7073 KB"/></supplementary-material><supplementary-material id="app2"><label>Checklist 1</label><p>PRISMA 2020 checklist.</p><media xlink:href="mental_v13i1e93307_app2.pdf" xlink:title="PDF File, 302 KB"/></supplementary-material></app-group></back></article>