<?xml version="1.0" encoding="UTF-8"?><!DOCTYPE article  PUBLIC "-//NLM//DTD Journal Publishing DTD v3.0 20080202//EN" "http://dtd.nlm.nih.gov/publishing/3.0/journalpublishing3.dtd"><article xmlns:mml="http://www.w3.org/1998/Math/MathML" xmlns:xlink="http://www.w3.org/1999/xlink" dtd-version="3.0" xml:lang="en" article-type="research article"><front><journal-meta><journal-id journal-id-type="publisher-id">OJBM</journal-id><journal-title-group><journal-title>Open Journal of Business and Management</journal-title></journal-title-group><issn pub-type="epub">2329-3284</issn><publisher><publisher-name>Scientific Research Publishing</publisher-name></publisher></journal-meta><article-meta><article-id pub-id-type="doi">10.4236/ojbm.2022.102047</article-id><article-id pub-id-type="publisher-id">OJBM-116054</article-id><article-categories><subj-group subj-group-type="heading"><subject>Articles</subject></subj-group><subj-group subj-group-type="Discipline-v2"><subject>Business&amp;Economics</subject></subj-group></article-categories><title-group><article-title>
 
 
  Measuring the Adequacy of Loss Distribution for the Ghanaian Auto Insurance Risk Exposure through Maximum Likelihood Estimation
 
</article-title></title-group><contrib-group><contrib contrib-type="author" xlink:type="simple"><name name-style="western"><surname>Jacob</surname><given-names>Azaare</given-names></name><xref ref-type="aff" rid="aff1"><sup>1</sup></xref></contrib><contrib contrib-type="author" xlink:type="simple"><name name-style="western"><surname>Zhao</surname><given-names>Wu</given-names></name><xref ref-type="aff" rid="aff1"><sup>1</sup></xref></contrib><contrib contrib-type="author" xlink:type="simple"><name name-style="western"><surname>Yingying</surname><given-names>Zhu</given-names></name><xref ref-type="aff" rid="aff2"><sup>2</sup></xref></contrib><contrib contrib-type="author" xlink:type="simple"><name name-style="western"><surname>Gabriel</surname><given-names>Armah</given-names></name><xref ref-type="aff" rid="aff3"><sup>3</sup></xref></contrib><contrib contrib-type="author" xlink:type="simple"><name name-style="western"><surname>Gideon</surname><given-names>Mensah Engmann</given-names></name><xref ref-type="aff" rid="aff4"><sup>4</sup></xref></contrib><contrib contrib-type="author" xlink:type="simple"><name name-style="western"><surname>Socrates</surname><given-names>Modzi Kwadwo</given-names></name><xref ref-type="aff" rid="aff5"><sup>5</sup></xref></contrib><contrib contrib-type="author" xlink:type="simple"><name name-style="western"><surname>Bright</surname><given-names>Nana Kwame Ahia</given-names></name><xref ref-type="aff" rid="aff1"><sup>1</sup></xref></contrib><contrib contrib-type="author" xlink:type="simple"><name name-style="western"><surname>Enock</surname><given-names>Mintah Ampaw</given-names></name><xref ref-type="aff" rid="aff6"><sup>6</sup></xref></contrib></contrib-group><aff id="aff6"><addr-line>Applied Mathematics Department, Faculty of Applied Science and Technology, Koforidua Technical University, Koforidua, Ghana</addr-line></aff><aff id="aff2"><addr-line>Business School of Chengdu, University of China, Chengdu, China</addr-line></aff><aff id="aff3"><addr-line>Department of Business Computing, School of Computing and Information Sciences, C.K Tedam University of Technology and Applied Sciences, Navrongo, Ghana</addr-line></aff><aff id="aff1"><addr-line>School of Management and Economics, University of Electronic Science and Technology of China, Chengdu, China</addr-line></aff><aff id="aff4"><addr-line>Department of Biometry, School of Mathematical Sciences, C.K Tedam University of Technology and Applied Sciences, Navrongo, Ghana</addr-line></aff><aff id="aff5"><addr-line>Faculty of Business, Economics and Social-Sciences, University of Hamburg, Hamburg, Germany</addr-line></aff><pub-date pub-type="epub"><day>11</day><month>02</month><year>2022</year></pub-date><volume>10</volume><issue>02</issue><fpage>846</fpage><lpage>859</lpage><history><date date-type="received"><day>14,</day>	<month>January</month>	<year>2022</year></date><date date-type="rev-recd"><day>19,</day>	<month>March</month>	<year>2022</year>	</date><date date-type="accepted"><day>22,</day>	<month>March</month>	<year>2022</year></date></history><permissions><copyright-statement>&#169; Copyright  2014 by authors and Scientific Research Publishing Inc. </copyright-statement><copyright-year>2014</copyright-year><license><license-p>This work is licensed under the Creative Commons Attribution International License (CC BY). http://creativecommons.org/licenses/by/4.0/</license-p></license></permissions><abstract><p>
 
 
  Loss distribution plays an influential role in evaluating risks from policy
  holders’ claims. Nevertheless, the auto insurance market in Ghana pays little attention to policyholders’ claims distribution, resulting in the market’s inefficiency. This study investigates the type of loss distribution function that best approximates policyholders’ claims in Ghana. We applied the Kullback-Leibler divergence, Kolmogorov Smirnov, Anderson-Darling statistical tests and maximum likelihood estimation (MLE) to estimate policyholders’ claims. The results suggest that Ghana’s auto policyholder’s claims are better approximated using the lognormal probability distribution. Through the lognormal distribution, the industry can adequately evaluate policyholders’ claims to minimize potential loss. Additionally, this distribution could enable the market
   
  reach decisions on premiums and expected profits theoretically.
 
</p></abstract><kwd-group><kwd>Leptokurtic</kwd><kwd> Loss Distribution</kwd><kwd> Policyholders Claims</kwd><kwd> Auto Insurance</kwd><kwd> Lognormal</kwd></kwd-group></article-meta></front><body><sec id="s1"><title>1. Introduction</title><p>The significant contribution of auto insurance markets in every country’s economic growth and development cannot be underrated (Azaare &amp; Wu, 2020; Azaare et al., 2021). Hence, safeguarding the market to understand insurers’ exposed risk probability distribution is essential (Azaare &amp; Wu, 2020; Stojakovic &amp; Jeremic, 2016). Insurers’ risk exposure may be defined as their susceptibility to potential losses. Every insurer has a duty to price premiums profitably, and this can be guided if the insurer has a clue on which probability law better approximates the risk posed by the policyholders (Ibiwoye et al., 2011; Packov&#225; &amp; Brebera, 2015; Walhin &amp; Paris, 1997). This, according to researchers (Brouhns et al., 2003; Azaare &amp; Wu, 2020; Pinquet, 2000; Spindler et al., 2014) guide the insurer in making proper evaluations and predictions to avoid or minimize the potential losses. From the above, it’s evident that there have been many uncertainties on the part of insurance companies concerning their risk exposure. The uncertainties have an adverse effect on these insurers and the economy they operate as a whole. The interest of loss distributions by authors in the area of insurance and finance can be associated with their ability in predicting and pricing of premiums, in most instances relying on historical claims from policyholders (Ahmad et al., 2020; Bali &amp; Theodossiou, 2008; Brouhns et al., 2003; Azaare &amp; Wu, 2020; Pinquet, 2000; Spindler et al., 2014; Tremblay, 1992).</p><p>Unlike developed countries where auto insurers have robust pricing systems that capture policyholders’ claims, the Ghana National Insurance Commission (NIC) uses a pricing system that pays little attention to the importance of claims histories (Azaare &amp; Wu, 2020). However, drawing from the above literature, the danger posed by policyholders’ claims can never be undermined. Notwithstanding, the claims ratios from most market players in Ghana according to the NIC annual reports are below the internationally accepted standards. The claim ratio is calculated as the net claims incurred divided by the Net Earned Premiums. It is an influential ratio indicating the strength an insurer exercise in paying claims and to some extent, how well policyholders are treated. Through this ratio, policyholders can measure how much they receive in return for each Ghana cedi premium paid to their insurers. According to the NIC 2018 annual report, the claim ratio overall market for the past years has been low as already posited compared to the internationally acceptable benchmark, which falls between 40% and 60%. In 2018, for instance, the market average increased to 42% from 37% in 2017. Though the market average performance falls within the internationally accepted standard, the figures being recorded by the company under consideration have shown some slight deviations. This company in 2017 recorded a ratio of 43%, which increased to 64% in 2018. If care is not taken, this figure could increase further or may not fall within the market average since historic records from <xref ref-type="fig" rid="fig1">Figure 1</xref> show some remarkable variations. In 2016 for example, this company recorded a ratio of 46%, which is a consecutive decline from the two previous years of 55% and 64% according to the NIC.</p><p>Aside the from market claims ratio, another influential indicator of probability is the total expense ratio (Management Expense + Commission expense). This ratio is determined as a percentage of the Net Earned Premium with an internationally accepted ratio usually less than 40%. As this ratio becomes larger, it implies that the company is inefficiently discharging its duties, and this is more likely to impact its prompt payment of claims to policyholders. In the year under which the sample is considered for this research (2018), the market mean expense ratio was 99%, whereas that of the insurer under consideration was 100%. The market average for this ratio has, over the years, not been so good. This is because almost all the key players in the market have been performing below the international standard which, indicates that the market in general is not that efficient to guarantee policyholders’ claims payment. Therefore, we argue that for an insurer to be financially solvent and avoid eroding policyholders trust, the claims from policyholders’ must not be taken for granted. Hence, it’s imperative to properly evaluate and predict policyholder’s claims distribution to help offset the market’s inefficiencies. Therefore, to guide the insurer in making proper evaluations and predictions to avoid or minimize potential losses that could end up eroding trust and to attain financial solvency, this study seeks to investigate the type of loss distribution function that best approximate the policyholders’ claims in Ghana using real data from a major insurance company.</p><p>Many researchers using a generalized linear model (GLM), apply the method of maximum likelihood estimation (MLE), method of moments and Bayesian estimation in their quest to have optimal pricing models for auto insurance based on policyholder’s claims histories and a priori rating variables (Bolanc&#233; et al., 2007; Azaare &amp; Wu, 2020). As postulated by Bali and Theodossiou (2008) and Fang (2003), the basic problem requires selection of loss models for severity and frequencies of claim. Probability distribution functions are forward-looking because they are founded on actual data (Sarpong, 2019), and to estimate them perfectly, one does not necessarily need long historical time series data. Moreover, probability distribution functions have the merit of being fairly free of mathematical priors, have the capacity of adapting instantaneously to any change in data and also dealing accurately with any intrinsic risks in the data (Bali &amp; Theodossiou, 2008; Sarpong, 2019).</p><p>In the past decades, the applications of loss distribution in finance and insurance have been very predominant. It is known that most insurance and financial assets return are positive (Klugman, Panjer, &amp; Willmot, 2012), have heavier tails and unimodal hump-shaped with much higher kurtosis that requires the application of distributions such as the exponential, gamma, Weibull, lognormal, the inverse Gaussian rather than the standard normal distribution (Ahmad et al., 2020; Cooray &amp; Ananda, 2005; Godin, Mayoral, &amp; Morales, 2012; Lane, 2000). Using the normal distribution, “an insurance risk pricing model was developed based on a measure of risk distortion” (Wang, 2000). This model was later modified by Godin et al. (2012) using the normal inverse Gaussian to accommodate heavier skewed tails data. Packov&#225; and Brebera (2015), through actuarial modeling found that the claims size of third-party liability insurance better fits the Pareto distribution than its contenders like the Weibull and the Gamma after performing a series of statistical tests. Sarpong (2019) demonstrates that the lognormal distribution optimally models the seasonal volatilities existing between the American dollar and the Ghana cedi. “A semi-parametric approach based on the generalized method of moments (GMM) to tackle the specification situation concerning claim frequency distributions has been proposed” (Fang, 2003). Extensive research by Bali and Theodossiou (2008) “estimate the conditional and unconditional value at risk (VaR) thresholds, evaluate the performance of three extreme value distributions, generalized Pareto distribution (GPD), generalized extreme value distribution (GEV), and Box-Cox-GEV, and four skewed fat-tailed distributions, skewed generalized error distribution (SGED), skewed generalized t (SGT), the exponential generalized beta of the second kind (EGB2), and inverse hyperbolic sign (IHS)”.</p><p>According to Egan (2011), several probability distributions were fitted to the daily percentage returns regarding Standard and Poor’s 500 portfolios, and the optimum selection was found to be the student t-distribution. An empirical study comparing compound distributions of stock returns shows that monthly data returns were accurately approximated by normal distribution (Akgiray &amp; Booth, 1987). According to Tucker and Pond (1988), “the distribution of exchange returns satisfies normality because of their long-tailed and leptokurtic behavior”. In providing a clear description to insurance data, Eling (2012) utilized two available data sets in insurance and demonstrated competition between the skew-normal and the skew-student t distribution compared with other distributions. Bolanc&#233; et al. (2007) shows evidence of the bivariate claims model by applying the skew-normal and log-skew-normal distributions using Spanish auto insurance datasets. As a parametric alternative in modeling heavy-tailed data, Ahn et al. (2012) applied the log-phase type distribution. In fitting auto insurance claims, performance comparison was made using popular insurance data by Kazemi and Noorizadeh (2015).</p><p>From the research work done on finance and insurance data, it has been noticed that though their distributions are mostly not normal and asymmetric the quest of approximating such data with appropriate probability function has always been successful. Therefore, to guide the insurer in making proper evaluations and predictions to avoid or minimize potential losses that could end up eroding trust and to attain financial solvency, this study seeks to investigate the type of loss distribution function that best approximate the policyholders’ claims in Ghana using real data from a major insurance company.</p></sec><sec id="s2"><title>2. Methods</title><sec id="s2_1"><title>2.1. Data Collection/Source of Data</title><p>Information on risk exposure and premiums for (n = 23,434) vehicle insurance policyholders was obtained throughout 2018 from a leading Ghanaian insurance company (Azaare et al., 2021). Due to existing market competition, this leading insurer opted to be unanimous. Out of the total sample, 3,733 (15.9%) drivers had reported claims with an average claim of 9,547.60 Cedi with minimum and maximum being respectively 25.00 and 746,016.00 Ghana cedi. The policyholders’ claim is the variable of interest in this research since we aimed to investigate the appropriate probability function that best approximates it. This variable has an average driver age of 49 years, whiles the minimum and the maximum ages stood at 21 and 76 years, respectively.</p></sec><sec id="s2_2"><title>2.2. Loss Distributions</title><p>Here, we reviewed some loss distribution functions with the sole aim of testing the one that best fits the data on policyholders’ claims. The data’s density plots under consideration shown in <xref ref-type="fig" rid="fig2">Figure 2</xref> and <xref ref-type="fig" rid="fig4">Figure 4</xref> exhibited the distribution’s features; hence, the justification of their selection for investigation. Also, the selection process was influenced by the Cullen and Frey plot in <xref ref-type="fig" rid="fig3">Figure 3</xref> obtained through the Bootstrapping method. This graph produced positive values of both skewness and kurtosis, indicating a heavier tail distribution. Therefore, the most likely distributions with the above features to fit the data are; lognormal, gamma, exponential and Weibull. We provide the Mathematical functions that make this possible.</p><sec id="s2_2_1"><title>2.2.1. The Lognormal Distribution Function</title><p>A random positive variable N is claimed to be log normally distributed because x = ln ( N ) being a random variable is normally distributed. With parameters μ and σ , the outcome probability distribution function indicating the normal distribution of the lognormal random variable ln ( N ) equal</p><p>f ( N ; μ , σ ) = { 2 π σ N 2 π σ N e − ln ( N − μ ) 2 2 σ 2 ,   N ≤ 0 0 ,                                                                 N &lt; 0</p><p>The lognormal distribution has its expectation and variance respectively as</p><p>E ( N ) = e μ + σ 2 2 / 2 , V a r ( N ) = e 2 ( μ + σ 2 ) − e 2 μ + σ 2</p></sec><sec id="s2_2_2"><title>2.2.2. The Gamma Distribution Function</title><p>A random variable N being continuous is said to follow the Gamma distribution function having parameters α &gt; 0 ; β &gt; 0 if has a PDF given as</p><p>f ( N ) = { 1 α β α N α − 1 e − β N , N &gt; 0 0 , otherwise</p><p>This distribution function has expectation and variance respectfully as E ( N ) = β and V a r ( N ) = β .</p></sec><sec id="s2_2_3"><title>2.2.3. The Exponential Distribution Function</title><p>With location and scale parameter respectively as λ &gt; 0 ; α &gt; 0 , a random variable N follows an exponential distribution function if f ( N ) = 1 α e − ( N − λ α ) ,</p><p>N ≥ λ ; α &gt; 0 . This distribution function has its expectation and variance respectfully as E ( N ) = α and V a r ( N ) = α .</p></sec><sec id="s2_2_4"><title>2.2.4. The Weibull Distribution Function</title><p>A random variable N follows the Weibull distribution with respectively the scale, the shape parameters η &gt; 0 ; β &gt; 0 , if the probability distribution function of N having threshold parameter λ is</p><p>f ( N ; η , β ) = β η β ( N − λ ) β − 1 e [ − ( N − λ η ) β ] , N ≥ λ , η , β &gt; 0 , N</p><p>The Weibull distribution has its expectation and variance, respectively as;</p><p>E ( N ) = η Γ ( 1 + 1 β ) + λ ,   V a r ( N ) = η 2 ( Γ ( 1 + 2 β ) − Γ 2 ( 1 + 1 β ) ) .</p></sec></sec><sec id="s2_3"><title>2.3. Quantification of Information Lost through Kullback-Leibler Divergence and Model Selection Information Criteria</title><p>To ascertain the quantum of information lost in our data, we evaluate the entropy of the probability distribution. The probability distributions’ entropy of the data is</p><p>E = − ∑ i = 0 N p ( x i ) log p ( x i )</p><p>The Kullback-Leibler divergence ( D K L ) is obtained by modifying this equation to give the actual information missing for approximating one probability distribution by another. The D K L is;</p><p>D K L ( p | | q ) = ∑ i = 0 N p ( x i ) log p ( x i ) q ( x i ) .</p><p>where the real and the fitted probability distributions are p ( x i ) and q ( x i ) respectively. It’s always desirable for the D K L value of a particular probability distribution to be smaller. The smaller the value, the less information lost when such probability distribution is used to approximate the data.</p><p>In the model selection process, relative values of different statistical distributions are compared using information criteria for an observed data. The information criteria employed to compare the best distribution for the data under consideration are the Bayesian Information Criterion (BIC) and the Akaike Information Criterion (AIC). These criteria though helps in obtaining the optimum selection, adjust the fit of the model and factors parsimoniously by taking fundamental interest considering the different parameters involved. Hence, with a preference for the smaller values, the best distribution fit was obtained through these statistics based on the log-likelihood function calculated at the MLE. Readers may see for example, (Ahmad et al., 2020; G&#243;mez-D&#233;niz &amp; Calder&#237;n-Ojeda, 2018; Azaare &amp; Wu, 2020), for details.</p></sec><sec id="s2_4"><title>2.4. Maximum Likelihood Estimation</title><p>Let’s assume that a random variable N = [ x 1 , x 2 , ⋯ , x n ] T be a vector of n independent observations drawn from a population with PDF x = ln ( N ) , where θ = [ θ 1 , θ 2 , ⋯ , θ q ] T denotes a vector for q unknown parameters. The likelihood</p><p>function L ( θ ; N ) is defined as L ( θ ; N ) = ∏ i = 0 n g ( x i ; θ ) . The value of θ for the maximum likelihood estimate θ ^ = θ ^ ( N ) is the one that maximizes L ( θ ; N ) .</p></sec></sec><sec id="s3"><title>3. Results and Discussions</title><p>The purpose of this paper was to investigate the PDF that adequately approximates the risk exposure of auto insurers. In doing so, we first produced a depicting plot of the data’s non-parametric density. As observed in <xref ref-type="fig" rid="fig2">Figure 2</xref>, the data is positively skewed since the graph tails towards the right direction. We employed the bootstrapping approach in obtaining the Cullen and Frey plot. In <xref ref-type="fig" rid="fig3">Figure 3</xref>, it’s confirmed that there are many possible loss distributions to approximate the data. The summary statistics from this graph produced a skewness and kurtosis of 14.81 and 373.88, respectively. This indicates leptokurtic data since from the statistics, it is noticed that the data has a heavier right tail compared to the left, and hence, the possible loss distribution is positively skewed. The distributions in <xref ref-type="fig" rid="fig3">Figure 3</xref> includes: Normal, Exponential, Logistic, Beta, Lognormal, Gamma and Weibull. However, the only distributions that will be looked at to ascertain which one the data follow are the Lognormal, Gamma, Exponential and the Weibull since they are the positively skewed distributions as</p><p>shown in <xref ref-type="fig" rid="fig4">Figure 4</xref>. The theoretical QQ and PP plots for these positively skewed distributions are shown in <xref ref-type="fig" rid="fig5">Figure 5</xref>. From this figure, it is clearly shown which distribution is more likely to approximate the data. These plots as indicated is not favoring the exponential distribution as it deviates out of the distribution functions hypothesized. Therefore, the most expected PDF could be Lognormal, Gamma or Weibull as confirmed by the graph shown in <xref ref-type="fig" rid="fig6">Figure 6</xref>.</p><p>The optimal distribution has the smallest AIC and BIC values but with the highest log-likelihood statistic. It is observed from <xref ref-type="table" rid="table1">Table 1</xref> that the lognormal distribution received the smallest AIC and BIC values with a high log-likelihood</p><table-wrap id="table1" ><label><xref ref-type="table" rid="table1">Table 1</xref></label><caption><title> The Goodness-of-fit Criteria, Log likelihood and Kullback-Leibler divergence Statistics for insurers risk exposure</title></caption><table><tbody><thead><tr><th align="center" valign="middle" >Loss distribution</th><th align="center" valign="middle" >Log likelihood</th><th align="center" valign="middle" >AIC</th><th align="center" valign="middle" >BIC</th><th align="center" valign="middle" >Kullback-Leibler divergence</th></tr></thead><tr><td align="center" valign="middle" >Gamma</td><td align="center" valign="middle" >−7,530.687</td><td align="center" valign="middle" >15,065.37</td><td align="center" valign="middle" >15,077.82</td><td align="center" valign="middle" >−2.30</td></tr><tr><td align="center" valign="middle" >Exponential</td><td align="center" valign="middle" >−11,463.57</td><td align="center" valign="middle" >22,929.14</td><td align="center" valign="middle" >22,935.37</td><td align="center" valign="middle" >−1.50</td></tr><tr><td align="center" valign="middle" >Lognormal</td><td align="center" valign="middle" >−7,322.54</td><td align="center" valign="middle" >14,649.09</td><td align="center" valign="middle" >14,661.54</td><td align="center" valign="middle" >−2.50</td></tr><tr><td align="center" valign="middle" >Weibull</td><td align="center" valign="middle" >−7,456.07</td><td align="center" valign="middle" >14,916.13</td><td align="center" valign="middle" >14,928.58</td><td align="center" valign="middle" >−2.47</td></tr></tbody></table></table-wrap><p>Note: Smaller statistics are preferred for Akaike’s information criteria (AIC), Bayesian information criteria (BIC) and Kullback-Leibler divergence. Bigger statistic is preferred for the Log likelihood estimates.</p><p>statistic, and hence the data is more expected to be fitted by the lognormal distribution. To be sure that the data is lognormal, we relied on goodness-of-fit information. The results obtained from the analysis using the Kolmogorov Smirnov and Anderson-Darling statistics show that the data is better approximated with the lognormal distribution. Therefore, it is ascertained that policyholders’ claims is better approximated using the lognormal distribution with μ = − 2.439 and σ = 0.248 . As far as these information criteria and test statistics are concerned with insurance and financial loss distributions, our findings are in line with several studies (Ahn et al., 2012; Azaare &amp; Wu, 2020; G&#243;mez-D&#233;niz &amp; Calder&#237;n-Ojeda, 2018; Dutta and Jason, 2011).</p><p>To further confirm that the insurers’ risk exposure is approximated by the lognormal distribution, data based on this distribution was simulated. The simulated data has μ = − 2.440 and σ = 0.250 . We then compared the simulated data from the lognormal distribution with our actual data (insurance claims). It’s observed from <xref ref-type="fig" rid="fig7">Figure 7</xref> and <xref ref-type="fig" rid="fig8">Figure 8</xref> that, the plots for both the original and the simulated risk exposures have similar characteristics. Also shown in <xref ref-type="fig" rid="fig9">Figure 9</xref> is the empirical density plot of the simulated data that can be compared with the original plot in <xref ref-type="fig" rid="fig2">Figure 2</xref>. Thus, based on the graphs in Figures 7-9, we noticed that these sets of data are very comparable and having almost all the points fallen alone the fitted curves from the theoretical empirical and cumulative distributions of the QQ and the PP plots. Therefore, the optimal distribution for the data is lognormal. To determine which of these distributions fits the data with less information, the Kullback-Leibler divergence test was performed. From <xref ref-type="table" rid="table1">Table 1</xref>, the various probability distribution functions and their associated Kullback-Leibler divergence values are shown. Our observation here is that smaller amount of information is lost by approximating the data using the lognormal distribution.</p><p>Finally, it was very needful to ascertain whether these two datasets have the same probability distribution. Therefore, with continuity correction, Wilcoxon-signed rank test was performed. It was observed from the test result that there was a p-value of 0.102. Hence, at 5% significant level, the two datasets are identical. Therefore, we can conclude that auto insurance policyholders claim in Ghana is approximated by the lognormal probability distribution.</p></sec><sec id="s4"><title>4. Conclusion and Practical Implications</title><p>Though insurers have a duty to manage policyholder’s risk, nonetheless, they are profit-making companies. Therefore, their survival in every economy depends on how they can properly evaluate their risk exposures to attain financial solvency. Hence, predicting policyholders’ claims probability distribution function would maximize profit. This paper has adequately established that the risk policyholders posed to insurers follows lognormal probability distribution. Using the lognormal probability distribution to evaluate the potential losses or policyholders’ risk will end up maximizing insurers’ profit. From the AIC, BIC, log-likelihood statistics, and the Wilcoxon signed-rank test, it was observed that this distribution properly approximates both the simulated and the original data. It was also observed from the Kullback-Leibler divergence statistics that the amount of information loss using the lognormal distribution to fit the data is less compared to the contending distributions. Thus, the lognormal is the best to approximate the auto insurance policyholders’ claims.</p><p>Managerially, this paper is expected to help insurers properly evaluate and manage policyholder’s risk to profit from it. Evaluating and predicting insurers’ risk exposures would translate into attaining financial solvency through maximization of profit. The lognormal distribution provides insurers useful and tractable mathematical features of the risk posed by policyholders. Aside from the useful and tractable features with the lognormal distribution, it also provides information to insurers to reach decisions on premiums loadings, expected profits, and necessary reserves needed to ensure profit margins and the effects of deductibles and reinsurance. The skewed nature of the lognormal distribution would provide insurance companies with the best alternative in modeling their risk exposures. Hence, it is recommended that auto insurers risk (policyholders’ claims) in Ghana and other developing economies are predicted using the lognormal distribution. Furthermore, while the lognormal distribution is statistically proven to provide a good fit for our data, we suggest that future research could also look into other long tailed distributions and categorize policyholders’ claims based on the insurance type.</p></sec><sec id="s5"><title>Acknowledgements</title><p>This work was supported by the National Science Foundation of China (project No: 71871044).</p></sec><sec id="s6"><title>Conflicts of Interest</title><p>Authors declare no conflict of interest.</p></sec><sec id="s7"><title>Cite this paper</title><p>Azaare, J., Wu, Z., Zhu, Y. Y., Armah, G., Engmann, G. M., Kwadwo, S. M., Ahia, B. N. K., &amp; Ampaw, E. M. (2022). Measuring the Adequacy of Loss Distribution for the Ghanaian Auto Insurance Risk Exposure through Maximum Likelihood Estimation. Open Journal of Business and Management, 10, 846-859. https://doi.org/10.4236/ojbm.2022.102047</p></sec></body><back><ref-list><title>References</title><ref id="scirp.116054-ref1"><label>1</label><mixed-citation publication-type="other" xlink:type="simple">Ahmad, Z., Mahmoudi, E., Sanku, D., &amp; Saima, K. K. (2020). Modeling Vehicle Insurance Loss Data Using a New Member of T-X Family of Distributions. Journal of Statistical Theory and Applications, 19, 133-147. https://doi.org/10.2991/jsta.d.200421.001</mixed-citation></ref><ref id="scirp.116054-ref2"><label>2</label><mixed-citation publication-type="other" xlink:type="simple">Ahn, S., Joseph, H. T. K., &amp; Vaidyanathan, R. (2012). A New Class of Models for Heavy Tailed Distributions in Finance and Insurance Risk. Insurance: Mathematics and Economics, 51, 43-52. https://doi.org/10.1016/j.insmatheco.2012.02.002</mixed-citation></ref><ref id="scirp.116054-ref3"><label>3</label><mixed-citation publication-type="other" xlink:type="simple">Akgiray, V., &amp; Geoffrey, G. B. (1987). Compound Distribution Models of Stock Returns: An Empirical Comparison. Journal of Financial Research, 10, 269-280.  
https://doi.org/10.1111/j.1475-6803.1987.tb00497.x</mixed-citation></ref><ref id="scirp.116054-ref4"><label>4</label><mixed-citation publication-type="other" xlink:type="simple">Azaare, J., &amp; Wu, Z. (2020). An Alternative Pricing System through Bayesian Estimates and Method of Moments in a Bonus-Malus Framework for the Ghanaian Auto Insurance Market. Journal of Risk and Financial Management, 13, Article No. 143.  
https://doi.org/10.3390/jrfm13070143</mixed-citation></ref><ref id="scirp.116054-ref5"><label>5</label><mixed-citation publication-type="other" xlink:type="simple">Azaare, J., Wu, Z., Gumah, B., Ampaw, E. M., &amp; Modzi, S. K. (2021). Auto Insurance Premiums in Ghana: An Autoregressive Distributed Lag Model Approach to Risk Exposure Variables. Journal of Psychology in Africa, 31, 362-368.  
https://doi.org/10.1080/14330237.2021.1952668</mixed-citation></ref><ref id="scirp.116054-ref6"><label>6</label><mixed-citation publication-type="other" xlink:type="simple">Bali, T. G., &amp; Panayiotis, T. (2008). Risk Measurement Performance of Alternative Distribution Functions. Journal of Risk and Insurance, 75, 411-437. 
https://doi.org/10.1111/j.1539-6975.2008.00266.x</mixed-citation></ref><ref id="scirp.116054-ref7"><label>7</label><mixed-citation publication-type="other" xlink:type="simple">Bolancé, C., Denuit, M., Guillén, M., &amp; Lambert, P. (2007). Greatest Accuracy Credibility with Dynamic Heterogeneity: The Harvey-Fernandes Model. Belgian Actuarial Bulleting, 7, 14-18 https://orbiuliegebe/handle/2268/167579</mixed-citation></ref><ref id="scirp.116054-ref8"><label>8</label><mixed-citation publication-type="other" xlink:type="simple">Brouhns, J., Guillen, N., &amp; Denuit, M. (2003). Bonus-Malus Scales in Segmented Tariffs with Stochastic Migration between Segments. Journal of Risk and Insurance, 70, 577-599.  
https://doi.org/10.1046/j.0022-4367.2003.00066.x</mixed-citation></ref><ref id="scirp.116054-ref9"><label>9</label><mixed-citation publication-type="other" xlink:type="simple">Cooray, K., &amp; Malwane, M. A. A. (2005). Modeling Actuarial Data with a Composite Lognormal-Pareto Model. Scandinavian Actuarial Journal, 2005, 321-334.  
https://doi.org/10.1080/03461230510009763</mixed-citation></ref><ref id="scirp.116054-ref10"><label>10</label><mixed-citation publication-type="other" xlink:type="simple">Dutta, K., &amp; Jason, P. (2011). A Tale of Tails: An Empirical Analysis of Loss Distribution Models for Estimating Operational Risk Capital. SSRN Electronic Journal.</mixed-citation></ref><ref id="scirp.116054-ref11"><label>11</label><mixed-citation publication-type="other" xlink:type="simple">Egan, W. J. (2011). The Distribution of S&amp;P 500 Index Returns. SSRN Electronic Journal.</mixed-citation></ref><ref id="scirp.116054-ref12"><label>12</label><mixed-citation publication-type="other" xlink:type="simple">Eling, M. (2012). Fitting Insurance Claims to Skewed Distributions: Are the Skew-Normal and Skew-Student Good Models? Insurance: Mathematics and Economics, 51, 239-248.  
https://doi.org/10.1016/j.insmatheco.2012.04.001</mixed-citation></ref><ref id="scirp.116054-ref13"><label>13</label><mixed-citation publication-type="other" xlink:type="simple">Fang, Y. (2003). Semi-Parametric Specification Tests for Discrete Probability Models. Journal of Risk and Insurance, 70, 73-84. https://doi.org/10.1111/1539-6975.00048</mixed-citation></ref><ref id="scirp.116054-ref14"><label>14</label><mixed-citation publication-type="other" xlink:type="simple">Godin, F., Silvia M., &amp; Manuel, M. (2012). Contingent Claim Pricing Using a Normal Inverse Gaussian Probability Distortion Operator. Journal of Risk and Insurance, 79, 841-866. https://doi.org/10.1111/j.1539-6975.2011.01445.x</mixed-citation></ref><ref id="scirp.116054-ref15"><label>15</label><mixed-citation publication-type="other" xlink:type="simple">Gómez-Déniz, E., &amp; Enrique, C. O. (2018). Multivariate Credibility in Bonus-Malus Systems Distinguishing between Different Types of Claims. Risks, 6, Article No. 34.  
https://doi.org/10.3390/risks6020034</mixed-citation></ref><ref id="scirp.116054-ref16"><label>16</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Ibiwoye</surname><given-names> A.</given-names></name>,<name name-style="western"><surname> Adeleke</surname><given-names> I. A.</given-names></name>,<name name-style="western"><surname> &amp; Aduloju</surname><given-names> S. A. </given-names></name>,<etal>et al</etal>. (<year>2011</year>)<article-title>. Quest for Optimal Bonus-Malus in Automobile Insurance in Developing Economies: An Actuarial Perspective</article-title><source> International Business Research</source><volume> 4</volume>,<fpage> 74</fpage>-<lpage>83</lpage>.<pub-id pub-id-type="doi"></pub-id></mixed-citation></ref><ref id="scirp.116054-ref17"><label>17</label><mixed-citation publication-type="journal" xlink:type="simple"><name name-style="western"><surname>Kazemi</surname><given-names> R.</given-names></name>,<name name-style="western"><surname> &amp; Monireh</surname><given-names> N. </given-names></name>,<etal>et al</etal>. (<year>2015</year>)<article-title>. A Comparison between Skew-Logistic and Skew-Normal Distributions</article-title><source> Matematika</source><volume> 31</volume>,<fpage> 15</fpage>-<lpage>24</lpage>.<pub-id pub-id-type="doi"></pub-id></mixed-citation></ref><ref id="scirp.116054-ref18"><label>18</label><mixed-citation publication-type="other" xlink:type="simple">Klugman, S. A., Harry H. P., &amp; Gordon, E. W. (2012). Loss Models: From Data to Decisions (3rd ed.). John Wiley &amp; Sons.</mixed-citation></ref><ref id="scirp.116054-ref19"><label>19</label><mixed-citation publication-type="other" xlink:type="simple">Lane, M. N. (2000). Pricing Risk Transfer Transactions. ASTIN Bulletin, 30, 259-293.  
https://doi.org/10.2143/AST.30.2.504635</mixed-citation></ref><ref id="scirp.116054-ref20"><label>20</label><mixed-citation publication-type="other" xlink:type="simple">Packová, V., &amp; Brebera, D. (2015). Loss Distributions in Insurance Risk Management. Proceedings of the International Conference on Economics and Business Administration (pp. 17-22).</mixed-citation></ref><ref id="scirp.116054-ref21"><label>21</label><mixed-citation publication-type="book" xlink:type="simple">Pinquet, J. (2000). Experience Rating through Heterogeneous Models. In G. Dionne, Ed., Handbook of Insurance (pp. 459-500). Amsterdam Kluwer Academic Publishers.  
https://doi.org/10.1007/978-94-010-0642-2_14</mixed-citation></ref><ref id="scirp.116054-ref22"><label>22</label><mixed-citation publication-type="other" xlink:type="simple">Sarpong, S. (2019). Estimating the Probability Distribution of the Exchange Rate between Ghana Cedi and American Dollar. Journal of King Saud University-Science, 31, 177-183. https://doi.org/10.1016/j.jksus.2018.04.023</mixed-citation></ref><ref id="scirp.116054-ref23"><label>23</label><mixed-citation publication-type="other" xlink:type="simple">Spindler, M., Joachim, W., &amp; Steffen, H. (2014). Asymmetric Information in the Market for Automobile Insurance: Evidence from Germany. Journal of Risk and Insurance, 81, 781-801. https://doi.org/10.1111/j.1539-6975.2013.12006.x</mixed-citation></ref><ref id="scirp.116054-ref24"><label>24</label><mixed-citation publication-type="other" xlink:type="simple">Stojakovic, A., &amp; Ljiljana, J. (2016). Development of the Insurance Sector and Economic Growth in Countries in Transition. Megatrend Revija, 13, 83-106. 
https://doi.org/10.5937/MegRev1603083S</mixed-citation></ref><ref id="scirp.116054-ref25"><label>25</label><mixed-citation publication-type="other" xlink:type="simple">Tremblay, L. (1992). Using the Poisson Inverse Gaussian in Bonus-Malus Systems. ASTIN Bulletin, 22, 97-106. https://doi.org/10.2143/AST.22.1.2005129</mixed-citation></ref><ref id="scirp.116054-ref26"><label>26</label><mixed-citation publication-type="other" xlink:type="simple">Tucker, A. L., &amp; Lallon, P. (1988). The Probability Distribution of Foreign Exchange Price Changes: Tests of Candidate Processes. The Review of Economics and Statistics, 70, 638-647. https://doi.org/10.2307/1935827</mixed-citation></ref><ref id="scirp.116054-ref27"><label>27</label><mixed-citation publication-type="other" xlink:type="simple">Walhin, J. F., &amp; Paris, J. (1997). Using Mixed Poisson Processes in Connection with Bonus-Malus Systems. Astin Bulletin, 29, 81-99. https://doi.org/10.2143/AST.29.1.504607</mixed-citation></ref><ref id="scirp.116054-ref28"><label>28</label><mixed-citation publication-type="other" xlink:type="simple">Wang, S. S. (2000). A Class of Distortion Operators for Pricing Financial and Insurance Risks. The Journal of Risk and Insurance, 67, 15-36. https://doi.org/10.2307/253675</mixed-citation></ref></ref-list></back></article>