diff mayachemtools/docs/scripts/html/code/MACCSKeysFingerprints.html @ 0:73ae111cf86f draft

Uploaded
author deepakjadmin
date Wed, 20 Jan 2016 11:55:01 -0500
parents
children
line wrap: on
line diff
--- /dev/null	Thu Jan 01 00:00:00 1970 +0000
+++ b/mayachemtools/docs/scripts/html/code/MACCSKeysFingerprints.html	Wed Jan 20 11:55:01 2016 -0500
@@ -0,0 +1,687 @@
+<html>
+<head>
+<title>MayaChemTools:Code:MACCSKeysFingerprints.pl</title>
+<meta http-equiv="content-type" content="text/html;charset=utf-8">
+<link rel="stylesheet" type="text/css" href="../../../css/MayaChemToolsCode.css">
+</head>
+<body leftmargin="20" rightmargin="20" topmargin="10" bottommargin="10">
+<br/>
+<center>
+<a href="http://www.mayachemtools.org" title="MayaChemTools Home"><img src="../../../images/MayaChemToolsLogo.gif" border="0" alt="MayaChemTools"></a>
+</center>
+<br/>
+<pre>
+   1 #!/usr/bin/perl -w
+   2 <span class="c">#</span>
+   3 <span class="c"># $RCSfile: MACCSKeysFingerprints.pl,v $</span>
+   4 <span class="c"># $Date: 2015/02/28 20:46:20 $</span>
+   5 <span class="c"># $Revision: 1.31 $</span>
+   6 <span class="c">#</span>
+   7 <span class="c"># Author: Manish Sud &lt;msud@san.rr.com&gt;</span>
+   8 <span class="c">#</span>
+   9 <span class="c"># Copyright (C) 2015 Manish Sud. All rights reserved.</span>
+  10 <span class="c">#</span>
+  11 <span class="c"># This file is part of MayaChemTools.</span>
+  12 <span class="c">#</span>
+  13 <span class="c"># MayaChemTools is free software; you can redistribute it and/or modify it under</span>
+  14 <span class="c"># the terms of the GNU Lesser General Public License as published by the Free</span>
+  15 <span class="c"># Software Foundation; either version 3 of the License, or (at your option) any</span>
+  16 <span class="c"># later version.</span>
+  17 <span class="c">#</span>
+  18 <span class="c"># MayaChemTools is distributed in the hope that it will be useful, but without</span>
+  19 <span class="c"># any warranty; without even the implied warranty of merchantability of fitness</span>
+  20 <span class="c"># for a particular purpose.  See the GNU Lesser General Public License for more</span>
+  21 <span class="c"># details.</span>
+  22 <span class="c">#</span>
+  23 <span class="c"># You should have received a copy of the GNU Lesser General Public License</span>
+  24 <span class="c"># along with MayaChemTools; if not, see &lt;http://www.gnu.org/licenses/&gt; or</span>
+  25 <span class="c"># write to the Free Software Foundation Inc., 59 Temple Place, Suite 330,</span>
+  26 <span class="c"># Boston, MA, 02111-1307, USA.</span>
+  27 <span class="c">#</span>
+  28 
+  29 <span class="k">use</span> <span class="w">strict</span><span class="sc">;</span>
+  30 <span class="k">use</span> <span class="w">FindBin</span><span class="sc">;</span> <span class="k">use</span> <span class="w">lib</span> <span class="q">&quot;$FindBin::Bin/../lib&quot;</span><span class="sc">;</span>
+  31 <span class="k">use</span> <span class="w">Getopt::Long</span><span class="sc">;</span>
+  32 <span class="k">use</span> <span class="w">File::Basename</span><span class="sc">;</span>
+  33 <span class="k">use</span> <span class="w">Text::ParseWords</span><span class="sc">;</span>
+  34 <span class="k">use</span> <span class="w">Benchmark</span><span class="sc">;</span>
+  35 <span class="k">use</span> <span class="w">FileUtil</span><span class="sc">;</span>
+  36 <span class="k">use</span> <span class="w">TextUtil</span><span class="sc">;</span>
+  37 <span class="k">use</span> <span class="w">SDFileUtil</span><span class="sc">;</span>
+  38 <span class="k">use</span> <span class="w">MoleculeFileIO</span><span class="sc">;</span>
+  39 <span class="k">use</span> <span class="w">FileIO::FingerprintsSDFileIO</span><span class="sc">;</span>
+  40 <span class="k">use</span> <span class="w">FileIO::FingerprintsTextFileIO</span><span class="sc">;</span>
+  41 <span class="k">use</span> <span class="w">FileIO::FingerprintsFPFileIO</span><span class="sc">;</span>
+  42 <span class="k">use</span> <span class="w">Fingerprints::MACCSKeys</span><span class="sc">;</span>
+  43 
+  44 <span class="k">my</span><span class="s">(</span><span class="i">$ScriptName</span><span class="cm">,</span> <span class="i">%Options</span><span class="cm">,</span> <span class="i">$StartTime</span><span class="cm">,</span> <span class="i">$EndTime</span><span class="cm">,</span> <span class="i">$TotalTime</span><span class="s">)</span><span class="sc">;</span>
+  45 
+  46 <span class="c"># Autoflush STDOUT</span>
+  47 <span class="i">$|</span> = <span class="n">1</span><span class="sc">;</span>
+  48 
+  49 <span class="c"># Starting message...</span>
+  50 <span class="i">$ScriptName</span> = <span class="i">basename</span><span class="s">(</span><span class="i">$0</span><span class="s">)</span><span class="sc">;</span>
+  51 <span class="k">print</span> <span class="q">&quot;\n$ScriptName: Starting...\n\n&quot;</span><span class="sc">;</span>
+  52 <span class="i">$StartTime</span> = <span class="w">new</span> <span class="w">Benchmark</span><span class="sc">;</span>
+  53 
+  54 <span class="c"># Get the options and setup script...</span>
+  55 <span class="i">SetupScriptUsage</span><span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+  56 <span class="k">if</span> <span class="s">(</span><span class="i">$Options</span>{<span class="w">help</span>} || <span class="i">@ARGV</span> &lt; <span class="n">1</span><span class="s">)</span> <span class="s">{</span>
+  57   <span class="k">die</span> <span class="i">GetUsageFromPod</span><span class="s">(</span><span class="q">&quot;$FindBin::Bin/$ScriptName&quot;</span><span class="s">)</span><span class="sc">;</span>
+  58 <span class="s">}</span>
+  59 
+  60 <span class="k">my</span><span class="s">(</span><span class="i">@SDFilesList</span><span class="s">)</span><span class="sc">;</span>
+  61 <span class="i">@SDFilesList</span> = <span class="i">ExpandFileNames</span><span class="s">(</span>\<span class="i">@ARGV</span><span class="cm">,</span> <span class="q">&quot;sdf sd&quot;</span><span class="s">)</span><span class="sc">;</span>
+  62 
+  63 <span class="c"># Process options...</span>
+  64 <span class="k">print</span> <span class="q">&quot;Processing options...\n&quot;</span><span class="sc">;</span>
+  65 <span class="k">my</span><span class="s">(</span><span class="i">%OptionsInfo</span><span class="s">)</span><span class="sc">;</span>
+  66 <span class="i">ProcessOptions</span><span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+  67 
+  68 <span class="c"># Setup information about input files...</span>
+  69 <span class="k">print</span> <span class="q">&quot;Checking input SD file(s)...\n&quot;</span><span class="sc">;</span>
+  70 <span class="k">my</span><span class="s">(</span><span class="i">%SDFilesInfo</span><span class="s">)</span><span class="sc">;</span>
+  71 <span class="i">RetrieveSDFilesInfo</span><span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+  72 
+  73 <span class="c"># Process input files..</span>
+  74 <span class="k">my</span><span class="s">(</span><span class="i">$FileIndex</span><span class="s">)</span><span class="sc">;</span>
+  75 <span class="k">if</span> <span class="s">(</span><span class="i">@SDFilesList</span> &gt; <span class="n">1</span><span class="s">)</span> <span class="s">{</span>
+  76   <span class="k">print</span> <span class="q">&quot;\nProcessing SD files...\n&quot;</span><span class="sc">;</span>
+  77 <span class="s">}</span>
+  78 <span class="k">for</span> <span class="i">$FileIndex</span> <span class="s">(</span><span class="n">0</span> .. <span class="i">$#SDFilesList</span><span class="s">)</span> <span class="s">{</span>
+  79   <span class="k">if</span> <span class="s">(</span><span class="i">$SDFilesInfo</span>{<span class="w">FileOkay</span>}[<span class="i">$FileIndex</span>]<span class="s">)</span> <span class="s">{</span>
+  80     <span class="k">print</span> <span class="q">&quot;\nProcessing file $SDFilesList[$FileIndex]...\n&quot;</span><span class="sc">;</span>
+  81     <span class="i">GenerateMACCSKeysFingerprints</span><span class="s">(</span><span class="i">$FileIndex</span><span class="s">)</span><span class="sc">;</span>
+  82   <span class="s">}</span>
+  83 <span class="s">}</span>
+  84 <span class="k">print</span> <span class="q">&quot;\n$ScriptName:Done...\n\n&quot;</span><span class="sc">;</span>
+  85 
+  86 <span class="i">$EndTime</span> = <span class="w">new</span> <span class="w">Benchmark</span><span class="sc">;</span>
+  87 <span class="i">$TotalTime</span> = <span class="w">timediff</span> <span class="s">(</span><span class="i">$EndTime</span><span class="cm">,</span> <span class="i">$StartTime</span><span class="s">)</span><span class="sc">;</span>
+  88 <span class="k">print</span> <span class="q">&quot;Total time: &quot;</span><span class="cm">,</span> <span class="i">timestr</span><span class="s">(</span><span class="i">$TotalTime</span><span class="s">)</span><span class="cm">,</span> <span class="q">&quot;\n&quot;</span><span class="sc">;</span>
+  89 
+  90 <span class="c">###############################################################################</span>
+  91 
+  92 <span class="c"># Generate fingerprints for a SD file...</span>
+  93 <span class="c">#</span>
+<a name="GenerateMACCSKeysFingerprints-"></a>  94 <span class="k">sub </span><span class="m">GenerateMACCSKeysFingerprints</span> <span class="s">{</span>
+  95   <span class="k">my</span><span class="s">(</span><span class="i">$FileIndex</span><span class="s">)</span> = <span class="i">@_</span><span class="sc">;</span>
+  96   <span class="k">my</span><span class="s">(</span><span class="i">$CmpdCount</span><span class="cm">,</span> <span class="i">$IgnoredCmpdCount</span><span class="cm">,</span> <span class="i">$SDFile</span><span class="cm">,</span> <span class="i">$MoleculeFileIO</span><span class="cm">,</span> <span class="i">$Molecule</span><span class="cm">,</span> <span class="i">$MACCSKeysFingerprints</span><span class="cm">,</span> <span class="i">$NewFPSDFileIO</span><span class="cm">,</span> <span class="i">$NewFPTextFileIO</span><span class="cm">,</span> <span class="i">$NewFPFileIO</span><span class="s">)</span><span class="sc">;</span>
+  97 
+  98   <span class="i">$SDFile</span> = <span class="i">$SDFilesList</span>[<span class="i">$FileIndex</span>]<span class="sc">;</span>
+  99 
+ 100   <span class="c"># Setup output files...</span>
+ 101   <span class="c">#</span>
+ 102   <span class="s">(</span><span class="i">$NewFPSDFileIO</span><span class="cm">,</span> <span class="i">$NewFPTextFileIO</span><span class="cm">,</span> <span class="i">$NewFPFileIO</span><span class="s">)</span> = <span class="i">SetupAndOpenOutputFiles</span><span class="s">(</span><span class="i">$FileIndex</span><span class="s">)</span><span class="sc">;</span>
+ 103 
+ 104   <span class="i">$MoleculeFileIO</span> = <span class="w">new</span> <span class="i">MoleculeFileIO</span><span class="s">(</span><span class="q">&#39;Name&#39;</span> <span class="cm">=&gt;</span> <span class="i">$SDFile</span><span class="s">)</span><span class="sc">;</span>
+ 105   <span class="i">$MoleculeFileIO</span><span class="i">-&gt;Open</span><span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 106 
+ 107   <span class="i">$CmpdCount</span> = <span class="n">0</span><span class="sc">;</span>
+ 108   <span class="i">$IgnoredCmpdCount</span> = <span class="n">0</span><span class="sc">;</span>
+ 109 
+ 110   <span class="j">COMPOUND:</span> <span class="k">while</span> <span class="s">(</span><span class="i">$Molecule</span> = <span class="i">$MoleculeFileIO</span><span class="i">-&gt;ReadMolecule</span><span class="s">(</span><span class="s">)</span><span class="s">)</span> <span class="s">{</span>
+ 111     <span class="i">$CmpdCount</span>++<span class="sc">;</span>
+ 112 
+ 113     <span class="c"># Filter compound data before calculating fingerprints...</span>
+ 114     <span class="k">if</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">Filter</span>}<span class="s">)</span> <span class="s">{</span>
+ 115       <span class="k">if</span> <span class="s">(</span><span class="i">CheckAndFilterCompound</span><span class="s">(</span><span class="i">$CmpdCount</span><span class="cm">,</span> <span class="i">$Molecule</span><span class="s">)</span><span class="s">)</span> <span class="s">{</span>
+ 116         <span class="i">$IgnoredCmpdCount</span>++<span class="sc">;</span>
+ 117         <span class="k">next</span> <span class="j">COMPOUND</span><span class="sc">;</span>
+ 118       <span class="s">}</span>
+ 119     <span class="s">}</span>
+ 120 
+ 121     <span class="i">$MACCSKeysFingerprints</span> = <span class="i">GenerateMoleculeFingerprints</span><span class="s">(</span><span class="i">$Molecule</span><span class="s">)</span><span class="sc">;</span>
+ 122     <span class="k">if</span> <span class="s">(</span>!<span class="i">$MACCSKeysFingerprints</span><span class="s">)</span> <span class="s">{</span>
+ 123       <span class="i">$IgnoredCmpdCount</span>++<span class="sc">;</span>
+ 124       <span class="i">ProcessIgnoredCompound</span><span class="s">(</span><span class="q">&#39;FingerprintsGenerationFailed&#39;</span><span class="cm">,</span> <span class="i">$CmpdCount</span><span class="cm">,</span> <span class="i">$Molecule</span><span class="s">)</span><span class="sc">;</span>
+ 125       <span class="k">next</span> <span class="j">COMPOUND</span><span class="sc">;</span>
+ 126     <span class="s">}</span>
+ 127 
+ 128     <span class="i">WriteDataToOutputFiles</span><span class="s">(</span><span class="i">$FileIndex</span><span class="cm">,</span> <span class="i">$CmpdCount</span><span class="cm">,</span> <span class="i">$Molecule</span><span class="cm">,</span> <span class="i">$MACCSKeysFingerprints</span><span class="cm">,</span> <span class="i">$NewFPSDFileIO</span><span class="cm">,</span> <span class="i">$NewFPTextFileIO</span><span class="cm">,</span> <span class="i">$NewFPFileIO</span><span class="s">)</span><span class="sc">;</span>
+ 129   <span class="s">}</span>
+ 130   <span class="i">$MoleculeFileIO</span><span class="i">-&gt;Close</span><span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 131 
+ 132   <span class="k">if</span> <span class="s">(</span><span class="i">$NewFPSDFileIO</span><span class="s">)</span> <span class="s">{</span>
+ 133     <span class="i">$NewFPSDFileIO</span><span class="i">-&gt;Close</span><span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 134   <span class="s">}</span>
+ 135   <span class="k">if</span> <span class="s">(</span><span class="i">$NewFPTextFileIO</span><span class="s">)</span> <span class="s">{</span>
+ 136     <span class="i">$NewFPTextFileIO</span><span class="i">-&gt;Close</span><span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 137   <span class="s">}</span>
+ 138   <span class="k">if</span> <span class="s">(</span><span class="i">$NewFPFileIO</span><span class="s">)</span> <span class="s">{</span>
+ 139     <span class="i">$NewFPFileIO</span><span class="i">-&gt;Close</span><span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 140   <span class="s">}</span>
+ 141 
+ 142   <span class="i">WriteFingerprintsGenerationSummaryStatistics</span><span class="s">(</span><span class="i">$CmpdCount</span><span class="cm">,</span> <span class="i">$IgnoredCmpdCount</span><span class="s">)</span><span class="sc">;</span>
+ 143 <span class="s">}</span>
+ 144 
+ 145 <span class="c"># Process compound being ignored due to problems in fingerprints geneation...</span>
+ 146 <span class="c">#</span>
+<a name="ProcessIgnoredCompound-"></a> 147 <span class="k">sub </span><span class="m">ProcessIgnoredCompound</span> <span class="s">{</span>
+ 148   <span class="k">my</span><span class="s">(</span><span class="i">$Mode</span><span class="cm">,</span> <span class="i">$CmpdCount</span><span class="cm">,</span> <span class="i">$Molecule</span><span class="s">)</span> = <span class="i">@_</span><span class="sc">;</span>
+ 149   <span class="k">my</span><span class="s">(</span><span class="i">$CmpdID</span><span class="cm">,</span> <span class="i">$DataFieldLabelAndValuesRef</span><span class="s">)</span><span class="sc">;</span>
+ 150 
+ 151   <span class="i">$DataFieldLabelAndValuesRef</span> = <span class="i">$Molecule</span><span class="i">-&gt;GetDataFieldLabelAndValues</span><span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 152   <span class="i">$CmpdID</span> = <span class="i">SetupCmpdIDForOutputFiles</span><span class="s">(</span><span class="i">$CmpdCount</span><span class="cm">,</span> <span class="i">$Molecule</span><span class="cm">,</span> <span class="i">$DataFieldLabelAndValuesRef</span><span class="s">)</span><span class="sc">;</span>
+ 153 
+ 154   <span class="j">MODE:</span> <span class="s">{</span>
+ 155     <span class="k">if</span> <span class="s">(</span><span class="i">$Mode</span> =~ <span class="q">/^ContainsNonElementalData$/i</span><span class="s">)</span> <span class="s">{</span>
+ 156       <span class="k">warn</span> <span class="q">&quot;\nWarning: Ignoring compound record number $CmpdCount with ID $CmpdID: Compound contains atom data corresponding to non-elemental atom symbol(s)...\n\n&quot;</span><span class="sc">;</span>
+ 157       <span class="k">next</span> <span class="j">MODE</span><span class="sc">;</span>
+ 158     <span class="s">}</span>
+ 159 
+ 160     <span class="k">if</span> <span class="s">(</span><span class="i">$Mode</span> =~ <span class="q">/^ContainsNoElementalData$/i</span><span class="s">)</span> <span class="s">{</span>
+ 161       <span class="k">warn</span> <span class="q">&quot;\nWarning: Ignoring compound record number $CmpdCount with ID $CmpdID: Compound contains no atom data...\n\n&quot;</span><span class="sc">;</span>
+ 162       <span class="k">next</span> <span class="j">MODE</span><span class="sc">;</span>
+ 163     <span class="s">}</span>
+ 164 
+ 165     <span class="k">if</span> <span class="s">(</span><span class="i">$Mode</span> =~ <span class="q">/^FingerprintsGenerationFailed$/i</span><span class="s">)</span> <span class="s">{</span>
+ 166       <span class="k">warn</span> <span class="q">&quot;\nWarning: Ignoring compound record number $CmpdCount with ID $CmpdID: Fingerprints generation didn&#39;t succeed...\n\n&quot;</span><span class="sc">;</span>
+ 167       <span class="k">next</span> <span class="j">MODE</span><span class="sc">;</span>
+ 168     <span class="s">}</span>
+ 169     <span class="k">warn</span> <span class="q">&quot;\nWarning: Ignoring compound record number $CmpdCount with ID $CmpdID: Fingerprints generation didn&#39;t succeed...\n\n&quot;</span><span class="sc">;</span>
+ 170   <span class="s">}</span>
+ 171 <span class="s">}</span>
+ 172 
+ 173 <span class="c"># Check and filter compounds....</span>
+ 174 <span class="c">#</span>
+<a name="CheckAndFilterCompound-"></a> 175 <span class="k">sub </span><span class="m">CheckAndFilterCompound</span> <span class="s">{</span>
+ 176   <span class="k">my</span><span class="s">(</span><span class="i">$CmpdCount</span><span class="cm">,</span> <span class="i">$Molecule</span><span class="s">)</span> = <span class="i">@_</span><span class="sc">;</span>
+ 177   <span class="k">my</span><span class="s">(</span><span class="i">$ElementCount</span><span class="cm">,</span> <span class="i">$NonElementCount</span><span class="s">)</span><span class="sc">;</span>
+ 178 
+ 179   <span class="s">(</span><span class="i">$ElementCount</span><span class="cm">,</span> <span class="i">$NonElementCount</span><span class="s">)</span> = <span class="i">$Molecule</span><span class="i">-&gt;GetNumOfElementsAndNonElements</span><span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 180 
+ 181   <span class="k">if</span> <span class="s">(</span><span class="i">$NonElementCount</span><span class="s">)</span> <span class="s">{</span>
+ 182     <span class="i">ProcessIgnoredCompound</span><span class="s">(</span><span class="q">&#39;ContainsNonElementalData&#39;</span><span class="cm">,</span> <span class="i">$CmpdCount</span><span class="cm">,</span> <span class="i">$Molecule</span><span class="s">)</span><span class="sc">;</span>
+ 183     <span class="k">return</span> <span class="n">1</span><span class="sc">;</span>
+ 184   <span class="s">}</span>
+ 185 
+ 186   <span class="k">if</span> <span class="s">(</span>!<span class="i">$ElementCount</span><span class="s">)</span> <span class="s">{</span>
+ 187     <span class="i">ProcessIgnoredCompound</span><span class="s">(</span><span class="q">&#39;ContainsNoElementalData&#39;</span><span class="cm">,</span> <span class="i">$CmpdCount</span><span class="cm">,</span> <span class="i">$Molecule</span><span class="s">)</span><span class="sc">;</span>
+ 188     <span class="k">return</span> <span class="n">1</span><span class="sc">;</span>
+ 189   <span class="s">}</span>
+ 190 
+ 191   <span class="k">return</span> <span class="n">0</span><span class="sc">;</span>
+ 192 <span class="s">}</span>
+ 193 
+ 194 <span class="c"># Write out compounds fingerprints generation summary statistics...</span>
+ 195 <span class="c">#</span>
+<a name="WriteFingerprintsGenerationSummaryStatistics-"></a> 196 <span class="k">sub </span><span class="m">WriteFingerprintsGenerationSummaryStatistics</span> <span class="s">{</span>
+ 197   <span class="k">my</span><span class="s">(</span><span class="i">$CmpdCount</span><span class="cm">,</span> <span class="i">$IgnoredCmpdCount</span><span class="s">)</span> = <span class="i">@_</span><span class="sc">;</span>
+ 198   <span class="k">my</span><span class="s">(</span><span class="i">$ProcessedCmpdCount</span><span class="s">)</span><span class="sc">;</span>
+ 199 
+ 200   <span class="i">$ProcessedCmpdCount</span> = <span class="i">$CmpdCount</span> - <span class="i">$IgnoredCmpdCount</span><span class="sc">;</span>
+ 201 
+ 202   <span class="k">print</span> <span class="q">&quot;\nNumber of compounds: $CmpdCount\n&quot;</span><span class="sc">;</span>
+ 203   <span class="k">print</span> <span class="q">&quot;Number of compounds processed successfully during fingerprints generation: $ProcessedCmpdCount\n&quot;</span><span class="sc">;</span>
+ 204   <span class="k">print</span> <span class="q">&quot;Number of compounds ignored during fingerprints generation: $IgnoredCmpdCount\n&quot;</span><span class="sc">;</span>
+ 205 <span class="s">}</span>
+ 206 
+ 207 <span class="c"># Open output files...</span>
+ 208 <span class="c">#</span>
+<a name="SetupAndOpenOutputFiles-"></a> 209 <span class="k">sub </span><span class="m">SetupAndOpenOutputFiles</span> <span class="s">{</span>
+ 210   <span class="k">my</span><span class="s">(</span><span class="i">$FileIndex</span><span class="s">)</span> = <span class="i">@_</span><span class="sc">;</span>
+ 211   <span class="k">my</span><span class="s">(</span><span class="i">$NewFPSDFile</span><span class="cm">,</span> <span class="i">$NewFPFile</span><span class="cm">,</span> <span class="i">$NewFPTextFile</span><span class="cm">,</span> <span class="i">$NewFPSDFileIO</span><span class="cm">,</span> <span class="i">$NewFPTextFileIO</span><span class="cm">,</span> <span class="i">$NewFPFileIO</span><span class="cm">,</span> <span class="i">%FingerprintsFileIOParams</span><span class="s">)</span><span class="sc">;</span>
+ 212 
+ 213   <span class="s">(</span><span class="i">$NewFPSDFileIO</span><span class="cm">,</span> <span class="i">$NewFPTextFileIO</span><span class="cm">,</span> <span class="i">$NewFPFileIO</span><span class="s">)</span> = <span class="s">(</span><span class="k">undef</span><span class="s">)</span> x <span class="n">3</span><span class="sc">;</span>
+ 214 
+ 215   <span class="c"># Setup common parameters for fingerprints file IO objects...</span>
+ 216   <span class="c">#</span>
+ 217   <span class="i">%FingerprintsFileIOParams</span> = <span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 218   <span class="k">if</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">Mode</span>} =~ <span class="q">/^MACCSKeyBits$/i</span><span class="s">)</span> <span class="s">{</span>
+ 219     <span class="i">%FingerprintsFileIOParams</span> = <span class="s">(</span><span class="q">&#39;Mode&#39;</span> <span class="cm">=&gt;</span> <span class="q">&#39;Write&#39;</span><span class="cm">,</span> <span class="q">&#39;Overwrite&#39;</span> <span class="cm">=&gt;</span> <span class="i">$OptionsInfo</span>{<span class="w">OverwriteFiles</span>}<span class="cm">,</span> <span class="q">&#39;FingerprintsStringMode&#39;</span> <span class="cm">=&gt;</span> <span class="q">&#39;FingerprintsBitVectorString&#39;</span><span class="cm">,</span> <span class="q">&#39;BitStringFormat&#39;</span> <span class="cm">=&gt;</span> <span class="i">$OptionsInfo</span>{<span class="w">BitStringFormat</span>}<span class="cm">,</span> <span class="q">&#39;BitsOrder&#39;</span> <span class="cm">=&gt;</span> <span class="i">$OptionsInfo</span>{<span class="w">BitsOrder</span>}<span class="s">)</span><span class="sc">;</span>
+ 220   <span class="s">}</span>
+ 221   <span class="k">elsif</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">Mode</span>} =~ <span class="q">/^MACCSKeyCount$/i</span><span class="s">)</span> <span class="s">{</span>
+ 222     <span class="i">%FingerprintsFileIOParams</span> = <span class="s">(</span><span class="q">&#39;Mode&#39;</span> <span class="cm">=&gt;</span> <span class="q">&#39;Write&#39;</span><span class="cm">,</span> <span class="q">&#39;Overwrite&#39;</span> <span class="cm">=&gt;</span> <span class="i">$OptionsInfo</span>{<span class="w">OverwriteFiles</span>}<span class="cm">,</span> <span class="q">&#39;FingerprintsStringMode&#39;</span> <span class="cm">=&gt;</span> <span class="q">&#39;FingerprintsVectorString&#39;</span><span class="cm">,</span> <span class="q">&#39;VectorStringFormat&#39;</span> <span class="cm">=&gt;</span> <span class="i">$OptionsInfo</span>{<span class="w">VectorStringFormat</span>}<span class="s">)</span><span class="sc">;</span>
+ 223   <span class="s">}</span>
+ 224 
+ 225   <span class="k">if</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">SDOutput</span>}<span class="s">)</span> <span class="s">{</span>
+ 226     <span class="i">$NewFPSDFile</span> = <span class="i">$SDFilesInfo</span>{<span class="w">SDOutFileNames</span>}[<span class="i">$FileIndex</span>]<span class="sc">;</span>
+ 227     <span class="k">print</span> <span class="q">&quot;Generating SD file $NewFPSDFile...\n&quot;</span><span class="sc">;</span>
+ 228     <span class="i">$NewFPSDFileIO</span> = <span class="w">new</span> <span class="i">FileIO::FingerprintsSDFileIO</span><span class="s">(</span><span class="q">&#39;Name&#39;</span> <span class="cm">=&gt;</span> <span class="i">$NewFPSDFile</span><span class="cm">,</span> <span class="i">%FingerprintsFileIOParams</span><span class="cm">,</span> <span class="q">&#39;FingerprintsFieldLabel&#39;</span> <span class="cm">=&gt;</span> <span class="i">$OptionsInfo</span>{<span class="w">FingerprintsLabel</span>}<span class="s">)</span><span class="sc">;</span>
+ 229     <span class="i">$NewFPSDFileIO</span><span class="i">-&gt;Open</span><span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 230   <span class="s">}</span>
+ 231 
+ 232   <span class="k">if</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">FPOutput</span>}<span class="s">)</span> <span class="s">{</span>
+ 233     <span class="i">$NewFPFile</span> = <span class="i">$SDFilesInfo</span>{<span class="w">FPOutFileNames</span>}[<span class="i">$FileIndex</span>]<span class="sc">;</span>
+ 234     <span class="k">print</span> <span class="q">&quot;Generating FP file $NewFPFile...\n&quot;</span><span class="sc">;</span>
+ 235     <span class="i">$NewFPFileIO</span> = <span class="w">new</span> <span class="i">FileIO::FingerprintsFPFileIO</span><span class="s">(</span><span class="q">&#39;Name&#39;</span> <span class="cm">=&gt;</span> <span class="i">$NewFPFile</span><span class="cm">,</span> <span class="i">%FingerprintsFileIOParams</span><span class="s">)</span><span class="sc">;</span>
+ 236     <span class="i">$NewFPFileIO</span><span class="i">-&gt;Open</span><span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 237   <span class="s">}</span>
+ 238 
+ 239   <span class="k">if</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">TextOutput</span>}<span class="s">)</span> <span class="s">{</span>
+ 240     <span class="k">my</span><span class="s">(</span><span class="i">$ColLabelsRef</span><span class="s">)</span><span class="sc">;</span>
+ 241 
+ 242     <span class="i">$NewFPTextFile</span> = <span class="i">$SDFilesInfo</span>{<span class="w">TextOutFileNames</span>}[<span class="i">$FileIndex</span>]<span class="sc">;</span>
+ 243     <span class="i">$ColLabelsRef</span> = <span class="i">SetupFPTextFileCoulmnLabels</span><span class="s">(</span><span class="i">$FileIndex</span><span class="s">)</span><span class="sc">;</span>
+ 244 
+ 245     <span class="k">print</span> <span class="q">&quot;Generating text file $NewFPTextFile...\n&quot;</span><span class="sc">;</span>
+ 246     <span class="i">$NewFPTextFileIO</span> = <span class="w">new</span> <span class="i">FileIO::FingerprintsTextFileIO</span><span class="s">(</span><span class="q">&#39;Name&#39;</span> <span class="cm">=&gt;</span> <span class="i">$NewFPTextFile</span><span class="cm">,</span> <span class="i">%FingerprintsFileIOParams</span><span class="cm">,</span> <span class="q">&#39;DataColLabels&#39;</span> <span class="cm">=&gt;</span> <span class="i">$ColLabelsRef</span><span class="cm">,</span> <span class="q">&#39;OutDelim&#39;</span> <span class="cm">=&gt;</span> <span class="i">$OptionsInfo</span>{<span class="w">OutDelim</span>}<span class="cm">,</span> <span class="q">&#39;OutQuote&#39;</span> <span class="cm">=&gt;</span> <span class="i">$OptionsInfo</span>{<span class="w">OutQuote</span>}<span class="s">)</span><span class="sc">;</span>
+ 247     <span class="i">$NewFPTextFileIO</span><span class="i">-&gt;Open</span><span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 248   <span class="s">}</span>
+ 249 
+ 250   <span class="k">return</span> <span class="s">(</span><span class="i">$NewFPSDFileIO</span><span class="cm">,</span> <span class="i">$NewFPTextFileIO</span><span class="cm">,</span> <span class="i">$NewFPFileIO</span><span class="s">)</span><span class="sc">;</span>
+ 251 <span class="s">}</span>
+ 252 
+ 253 <span class="c"># Write fingerpritns and other data to appropriate output files...</span>
+ 254 <span class="c">#</span>
+<a name="WriteDataToOutputFiles-"></a> 255 <span class="k">sub </span><span class="m">WriteDataToOutputFiles</span> <span class="s">{</span>
+ 256   <span class="k">my</span><span class="s">(</span><span class="i">$FileIndex</span><span class="cm">,</span> <span class="i">$CmpdCount</span><span class="cm">,</span> <span class="i">$Molecule</span><span class="cm">,</span> <span class="i">$MACCSKeysFingerprints</span><span class="cm">,</span> <span class="i">$NewFPSDFileIO</span><span class="cm">,</span> <span class="i">$NewFPTextFileIO</span><span class="cm">,</span> <span class="i">$NewFPFileIO</span><span class="s">)</span> = <span class="i">@_</span><span class="sc">;</span>
+ 257   <span class="k">my</span><span class="s">(</span><span class="i">$DataFieldLabelAndValuesRef</span><span class="s">)</span><span class="sc">;</span>
+ 258 
+ 259   <span class="i">$DataFieldLabelAndValuesRef</span> = <span class="k">undef</span><span class="sc">;</span>
+ 260   <span class="k">if</span> <span class="s">(</span><span class="i">$NewFPTextFileIO</span> || <span class="i">$NewFPFileIO</span><span class="s">)</span> <span class="s">{</span>
+ 261     <span class="i">$DataFieldLabelAndValuesRef</span> = <span class="i">$Molecule</span><span class="i">-&gt;GetDataFieldLabelAndValues</span><span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 262   <span class="s">}</span>
+ 263 
+ 264   <span class="k">if</span> <span class="s">(</span><span class="i">$NewFPSDFileIO</span><span class="s">)</span> <span class="s">{</span>
+ 265     <span class="k">my</span><span class="s">(</span><span class="i">$CmpdString</span><span class="s">)</span><span class="sc">;</span>
+ 266 
+ 267     <span class="i">$CmpdString</span> = <span class="i">$Molecule</span><span class="i">-&gt;GetInputMoleculeString</span><span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 268     <span class="i">$NewFPSDFileIO</span><span class="i">-&gt;WriteFingerprints</span><span class="s">(</span><span class="i">$MACCSKeysFingerprints</span><span class="cm">,</span> <span class="i">$CmpdString</span><span class="s">)</span><span class="sc">;</span>
+ 269   <span class="s">}</span>
+ 270 
+ 271   <span class="k">if</span> <span class="s">(</span><span class="i">$NewFPTextFileIO</span><span class="s">)</span> <span class="s">{</span>
+ 272     <span class="k">my</span><span class="s">(</span><span class="i">$ColValuesRef</span><span class="s">)</span><span class="sc">;</span>
+ 273 
+ 274     <span class="i">$ColValuesRef</span> = <span class="i">SetupFPTextFileCoulmnValues</span><span class="s">(</span><span class="i">$FileIndex</span><span class="cm">,</span> <span class="i">$CmpdCount</span><span class="cm">,</span> <span class="i">$Molecule</span><span class="cm">,</span> <span class="i">$DataFieldLabelAndValuesRef</span><span class="s">)</span><span class="sc">;</span>
+ 275     <span class="i">$NewFPTextFileIO</span><span class="i">-&gt;WriteFingerprints</span><span class="s">(</span><span class="i">$MACCSKeysFingerprints</span><span class="cm">,</span> <span class="i">$ColValuesRef</span><span class="s">)</span><span class="sc">;</span>
+ 276   <span class="s">}</span>
+ 277 
+ 278   <span class="k">if</span> <span class="s">(</span><span class="i">$NewFPFileIO</span><span class="s">)</span> <span class="s">{</span>
+ 279     <span class="k">my</span><span class="s">(</span><span class="i">$CompoundID</span><span class="s">)</span><span class="sc">;</span>
+ 280 
+ 281     <span class="i">$CompoundID</span> = <span class="i">SetupCmpdIDForOutputFiles</span><span class="s">(</span><span class="i">$CmpdCount</span><span class="cm">,</span> <span class="i">$Molecule</span><span class="cm">,</span> <span class="i">$DataFieldLabelAndValuesRef</span><span class="s">)</span><span class="sc">;</span>
+ 282     <span class="i">$NewFPFileIO</span><span class="i">-&gt;WriteFingerprints</span><span class="s">(</span><span class="i">$MACCSKeysFingerprints</span><span class="cm">,</span> <span class="i">$CompoundID</span><span class="s">)</span><span class="sc">;</span>
+ 283   <span class="s">}</span>
+ 284 <span class="s">}</span>
+ 285 
+ 286 <span class="c"># Generate approriate column labels for FPText output file...</span>
+ 287 <span class="c">#</span>
+<a name="SetupFPTextFileCoulmnLabels-"></a> 288 <span class="k">sub </span><span class="m">SetupFPTextFileCoulmnLabels</span> <span class="s">{</span>
+ 289   <span class="k">my</span><span class="s">(</span><span class="i">$FileIndex</span><span class="s">)</span> = <span class="i">@_</span><span class="sc">;</span>
+ 290   <span class="k">my</span><span class="s">(</span><span class="i">$Line</span><span class="cm">,</span> <span class="i">@ColLabels</span><span class="s">)</span><span class="sc">;</span>
+ 291 
+ 292   <span class="i">@ColLabels</span> = <span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 293   <span class="k">if</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">DataFieldsMode</span>} =~ <span class="q">/^All$/i</span><span class="s">)</span> <span class="s">{</span>
+ 294     <span class="k">push</span> <span class="i">@ColLabels</span><span class="cm">,</span> <span class="i">@</span>{<span class="i">$SDFilesInfo</span>{<span class="w">AllDataFieldsRef</span>}[<span class="i">$FileIndex</span>]}<span class="sc">;</span>
+ 295   <span class="s">}</span>
+ 296   <span class="k">elsif</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">DataFieldsMode</span>} =~ <span class="q">/^Common$/i</span><span class="s">)</span> <span class="s">{</span>
+ 297     <span class="k">push</span> <span class="i">@ColLabels</span><span class="cm">,</span> <span class="i">@</span>{<span class="i">$SDFilesInfo</span>{<span class="w">CommonDataFieldsRef</span>}[<span class="i">$FileIndex</span>]}<span class="sc">;</span>
+ 298   <span class="s">}</span>
+ 299   <span class="k">elsif</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">DataFieldsMode</span>} =~ <span class="q">/^Specify$/i</span><span class="s">)</span> <span class="s">{</span>
+ 300     <span class="k">push</span> <span class="i">@ColLabels</span><span class="cm">,</span> <span class="i">@</span>{<span class="i">$OptionsInfo</span>{<span class="w">SpecifiedDataFields</span>}}<span class="sc">;</span>
+ 301   <span class="s">}</span>
+ 302   <span class="k">elsif</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">DataFieldsMode</span>} =~ <span class="q">/^CompoundID$/i</span><span class="s">)</span> <span class="s">{</span>
+ 303     <span class="k">push</span> <span class="i">@ColLabels</span><span class="cm">,</span> <span class="i">$OptionsInfo</span>{<span class="w">CompoundIDLabel</span>}<span class="sc">;</span>
+ 304   <span class="s">}</span>
+ 305   <span class="c"># Add fingerprints label...</span>
+ 306   <span class="k">push</span> <span class="i">@ColLabels</span><span class="cm">,</span> <span class="i">$OptionsInfo</span>{<span class="w">FingerprintsLabel</span>}<span class="sc">;</span>
+ 307 
+ 308   <span class="k">return</span> \<span class="i">@ColLabels</span><span class="sc">;</span>
+ 309 <span class="s">}</span>
+ 310 
+ 311 <span class="c"># Generate column values FPText output file..</span>
+ 312 <span class="c">#</span>
+<a name="SetupFPTextFileCoulmnValues-"></a> 313 <span class="k">sub </span><span class="m">SetupFPTextFileCoulmnValues</span> <span class="s">{</span>
+ 314   <span class="k">my</span><span class="s">(</span><span class="i">$FileIndex</span><span class="cm">,</span> <span class="i">$CmpdCount</span><span class="cm">,</span> <span class="i">$Molecule</span><span class="cm">,</span> <span class="i">$DataFieldLabelAndValuesRef</span><span class="s">)</span> = <span class="i">@_</span><span class="sc">;</span>
+ 315   <span class="k">my</span><span class="s">(</span><span class="i">@ColValues</span><span class="s">)</span><span class="sc">;</span>
+ 316 
+ 317   <span class="i">@ColValues</span> = <span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 318   <span class="k">if</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">DataFieldsMode</span>} =~ <span class="q">/^CompoundID$/i</span><span class="s">)</span> <span class="s">{</span>
+ 319     <span class="k">push</span> <span class="i">@ColValues</span><span class="cm">,</span> <span class="i">SetupCmpdIDForOutputFiles</span><span class="s">(</span><span class="i">$CmpdCount</span><span class="cm">,</span> <span class="i">$Molecule</span><span class="cm">,</span> <span class="i">$DataFieldLabelAndValuesRef</span><span class="s">)</span><span class="sc">;</span>
+ 320   <span class="s">}</span>
+ 321   <span class="k">elsif</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">DataFieldsMode</span>} =~ <span class="q">/^All$/i</span><span class="s">)</span> <span class="s">{</span>
+ 322     <span class="i">@ColValues</span> = <span class="k">map</span> <span class="s">{</span> <span class="k">exists</span> <span class="i">$DataFieldLabelAndValuesRef</span>-&gt;{<span class="i">$_</span>} ? <span class="i">$DataFieldLabelAndValuesRef</span>-&gt;{<span class="i">$_</span>} <span class="co">:</span> <span class="q">&#39;&#39;</span><span class="s">}</span> <span class="i">@</span>{<span class="i">$SDFilesInfo</span>{<span class="w">AllDataFieldsRef</span>}[<span class="i">$FileIndex</span>]}<span class="sc">;</span>
+ 323   <span class="s">}</span>
+ 324   <span class="k">elsif</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">DataFieldsMode</span>} =~ <span class="q">/^Common$/i</span><span class="s">)</span> <span class="s">{</span>
+ 325     <span class="i">@ColValues</span> = <span class="k">map</span> <span class="s">{</span> <span class="k">exists</span> <span class="i">$DataFieldLabelAndValuesRef</span>-&gt;{<span class="i">$_</span>} ? <span class="i">$DataFieldLabelAndValuesRef</span>-&gt;{<span class="i">$_</span>} <span class="co">:</span> <span class="q">&#39;&#39;</span><span class="s">}</span> <span class="i">@</span>{<span class="i">$SDFilesInfo</span>{<span class="w">CommonDataFieldsRef</span>}[<span class="i">$FileIndex</span>]}<span class="sc">;</span>
+ 326   <span class="s">}</span>
+ 327   <span class="k">elsif</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">DataFieldsMode</span>} =~ <span class="q">/^Specify$/i</span><span class="s">)</span> <span class="s">{</span>
+ 328     <span class="i">@ColValues</span> = <span class="k">map</span> <span class="s">{</span> <span class="k">exists</span> <span class="i">$DataFieldLabelAndValuesRef</span>-&gt;{<span class="i">$_</span>} ? <span class="i">$DataFieldLabelAndValuesRef</span>-&gt;{<span class="i">$_</span>} <span class="co">:</span> <span class="q">&#39;&#39;</span><span class="s">}</span> <span class="i">@</span>{<span class="i">$OptionsInfo</span>{<span class="w">SpecifiedDataFields</span>}}<span class="sc">;</span>
+ 329   <span class="s">}</span>
+ 330 
+ 331   <span class="k">return</span> \<span class="i">@ColValues</span><span class="sc">;</span>
+ 332 <span class="s">}</span>
+ 333 
+ 334 <span class="c"># Generate compound ID for FP and FPText output files..</span>
+ 335 <span class="c">#</span>
+<a name="SetupCmpdIDForOutputFiles-"></a> 336 <span class="k">sub </span><span class="m">SetupCmpdIDForOutputFiles</span> <span class="s">{</span>
+ 337   <span class="k">my</span><span class="s">(</span><span class="i">$CmpdCount</span><span class="cm">,</span> <span class="i">$Molecule</span><span class="cm">,</span> <span class="i">$DataFieldLabelAndValuesRef</span><span class="s">)</span> = <span class="i">@_</span><span class="sc">;</span>
+ 338   <span class="k">my</span><span class="s">(</span><span class="i">$CmpdID</span><span class="s">)</span><span class="sc">;</span>
+ 339 
+ 340   <span class="i">$CmpdID</span> = <span class="q">&#39;&#39;</span><span class="sc">;</span>
+ 341   <span class="k">if</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">CompoundIDMode</span>} =~ <span class="q">/^MolNameOrLabelPrefix$/i</span><span class="s">)</span> <span class="s">{</span>
+ 342     <span class="k">my</span><span class="s">(</span><span class="i">$MolName</span><span class="s">)</span><span class="sc">;</span>
+ 343     <span class="i">$MolName</span> = <span class="i">$Molecule</span><span class="i">-&gt;GetName</span><span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 344     <span class="i">$CmpdID</span> = <span class="i">$MolName</span> ? <span class="i">$MolName</span> <span class="co">:</span> <span class="q">&quot;$OptionsInfo{CompoundID}${CmpdCount}&quot;</span><span class="sc">;</span>
+ 345   <span class="s">}</span>
+ 346   <span class="k">elsif</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">CompoundIDMode</span>} =~ <span class="q">/^LabelPrefix$/i</span><span class="s">)</span> <span class="s">{</span>
+ 347     <span class="i">$CmpdID</span> = <span class="q">&quot;$OptionsInfo{CompoundID}${CmpdCount}&quot;</span><span class="sc">;</span>
+ 348   <span class="s">}</span>
+ 349   <span class="k">elsif</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">CompoundIDMode</span>} =~ <span class="q">/^DataField$/i</span><span class="s">)</span> <span class="s">{</span>
+ 350     <span class="k">my</span><span class="s">(</span><span class="i">$SpecifiedDataField</span><span class="s">)</span><span class="sc">;</span>
+ 351     <span class="i">$SpecifiedDataField</span> = <span class="i">$OptionsInfo</span>{<span class="w">CompoundID</span>}<span class="sc">;</span>
+ 352     <span class="i">$CmpdID</span> = <span class="k">exists</span> <span class="i">$DataFieldLabelAndValuesRef</span>-&gt;{<span class="i">$SpecifiedDataField</span>} ? <span class="i">$DataFieldLabelAndValuesRef</span>-&gt;{<span class="i">$SpecifiedDataField</span>} <span class="co">:</span> <span class="q">&#39;&#39;</span><span class="sc">;</span>
+ 353   <span class="s">}</span>
+ 354   <span class="k">elsif</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">CompoundIDMode</span>} =~ <span class="q">/^MolName$/i</span><span class="s">)</span> <span class="s">{</span>
+ 355     <span class="i">$CmpdID</span> = <span class="i">$Molecule</span><span class="i">-&gt;GetName</span><span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 356   <span class="s">}</span>
+ 357   <span class="k">return</span> <span class="i">$CmpdID</span><span class="sc">;</span>
+ 358 <span class="s">}</span>
+ 359 
+ 360 <span class="c"># Generate fingerprints for molecule...</span>
+ 361 <span class="c">#</span>
+<a name="GenerateMoleculeFingerprints-"></a> 362 <span class="k">sub </span><span class="m">GenerateMoleculeFingerprints</span> <span class="s">{</span>
+ 363   <span class="k">my</span><span class="s">(</span><span class="i">$Molecule</span><span class="s">)</span> = <span class="i">@_</span><span class="sc">;</span>
+ 364   <span class="k">my</span><span class="s">(</span><span class="i">$MACCSKeysFingerprints</span><span class="s">)</span><span class="sc">;</span>
+ 365 
+ 366   <span class="k">if</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">KeepLargestComponent</span>}<span class="s">)</span> <span class="s">{</span>
+ 367     <span class="i">$Molecule</span><span class="i">-&gt;KeepLargestComponent</span><span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 368   <span class="s">}</span>
+ 369   <span class="k">if</span> <span class="s">(</span>!<span class="i">$Molecule</span><span class="i">-&gt;DetectRings</span><span class="s">(</span><span class="s">)</span><span class="s">)</span> <span class="s">{</span>
+ 370     <span class="k">return</span> <span class="k">undef</span><span class="sc">;</span>
+ 371   <span class="s">}</span>
+ 372   <span class="i">$Molecule</span><span class="i">-&gt;SetAromaticityModel</span><span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">AromaticityModel</span>}<span class="s">)</span><span class="sc">;</span>
+ 373   <span class="i">$Molecule</span><span class="i">-&gt;DetectAromaticity</span><span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 374 
+ 375   <span class="i">$MACCSKeysFingerprints</span> = <span class="k">undef</span><span class="sc">;</span>
+ 376   <span class="k">if</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">Mode</span>} =~ <span class="q">/^MACCSKeyBits$/i</span><span class="s">)</span> <span class="s">{</span>
+ 377     <span class="i">$MACCSKeysFingerprints</span> = <span class="w">new</span> <span class="i">Fingerprints::MACCSKeys</span><span class="s">(</span><span class="q">&#39;Molecule&#39;</span> <span class="cm">=&gt;</span> <span class="i">$Molecule</span><span class="cm">,</span> <span class="q">&#39;Type&#39;</span> <span class="cm">=&gt;</span> <span class="q">&#39;MACCSKeyBits&#39;</span><span class="cm">,</span> <span class="q">&#39;Size&#39;</span> <span class="cm">=&gt;</span> <span class="i">$OptionsInfo</span>{<span class="w">Size</span>}<span class="s">)</span><span class="sc">;</span>
+ 378   <span class="s">}</span>
+ 379   <span class="k">elsif</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">Mode</span>} =~ <span class="q">/^MACCSKeyCount$/i</span><span class="s">)</span> <span class="s">{</span>
+ 380     <span class="i">$MACCSKeysFingerprints</span> = <span class="w">new</span> <span class="i">Fingerprints::MACCSKeys</span><span class="s">(</span><span class="q">&#39;Molecule&#39;</span> <span class="cm">=&gt;</span> <span class="i">$Molecule</span><span class="cm">,</span> <span class="q">&#39;Type&#39;</span> <span class="cm">=&gt;</span> <span class="q">&#39;MACCSKeyCount&#39;</span><span class="cm">,</span> <span class="q">&#39;Size&#39;</span> <span class="cm">=&gt;</span> <span class="i">$OptionsInfo</span>{<span class="w">Size</span>}<span class="s">)</span><span class="sc">;</span>
+ 381   <span class="s">}</span>
+ 382   <span class="k">else</span> <span class="s">{</span>
+ 383     <span class="k">die</span> <span class="q">&quot;Error: The value specified, $Options{mode}, for option \&quot;-m, --mode\&quot; is not valid. Allowed values: MACCSKeyBits or MACCSKeyCount\n&quot;</span><span class="sc">;</span>
+ 384   <span class="s">}</span>
+ 385   <span class="i">$MACCSKeysFingerprints</span><span class="i">-&gt;GenerateMACCSKeys</span><span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 386 
+ 387   <span class="k">return</span> <span class="i">$MACCSKeysFingerprints</span><span class="sc">;</span>
+ 388 <span class="s">}</span>
+ 389 
+ 390 <span class="c"># Retrieve information about SD files...</span>
+ 391 <span class="c">#</span>
+<a name="RetrieveSDFilesInfo-"></a> 392 <span class="k">sub </span><span class="m">RetrieveSDFilesInfo</span> <span class="s">{</span>
+ 393   <span class="k">my</span><span class="s">(</span><span class="i">$SDFile</span><span class="cm">,</span> <span class="i">$Index</span><span class="cm">,</span> <span class="i">$FileDir</span><span class="cm">,</span> <span class="i">$FileExt</span><span class="cm">,</span> <span class="i">$FileName</span><span class="cm">,</span> <span class="i">$OutFileRoot</span><span class="cm">,</span> <span class="i">$TextOutFileExt</span><span class="cm">,</span> <span class="i">$SDOutFileExt</span><span class="cm">,</span> <span class="i">$FPOutFileExt</span><span class="cm">,</span> <span class="i">$NewSDFileName</span><span class="cm">,</span> <span class="i">$NewFPFileName</span><span class="cm">,</span> <span class="i">$NewTextFileName</span><span class="cm">,</span> <span class="i">$CheckDataField</span><span class="cm">,</span> <span class="i">$CollectDataFields</span><span class="cm">,</span> <span class="i">$AllDataFieldsRef</span><span class="cm">,</span> <span class="i">$CommonDataFieldsRef</span><span class="s">)</span><span class="sc">;</span>
+ 394 
+ 395   <span class="i">%SDFilesInfo</span> = <span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 396   <span class="i">@</span>{<span class="i">$SDFilesInfo</span>{<span class="w">FileOkay</span>}} = <span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 397   <span class="i">@</span>{<span class="i">$SDFilesInfo</span>{<span class="w">OutFileRoot</span>}} = <span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 398   <span class="i">@</span>{<span class="i">$SDFilesInfo</span>{<span class="w">SDOutFileNames</span>}} = <span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 399   <span class="i">@</span>{<span class="i">$SDFilesInfo</span>{<span class="w">FPOutFileNames</span>}} = <span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 400   <span class="i">@</span>{<span class="i">$SDFilesInfo</span>{<span class="w">TextOutFileNames</span>}} = <span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 401   <span class="i">@</span>{<span class="i">$SDFilesInfo</span>{<span class="w">AllDataFieldsRef</span>}} = <span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 402   <span class="i">@</span>{<span class="i">$SDFilesInfo</span>{<span class="w">CommonDataFieldsRef</span>}} = <span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 403 
+ 404   <span class="i">$CheckDataField</span> = <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">TextOutput</span>} &amp;&amp; <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">DataFieldsMode</span>} =~ <span class="q">/^CompoundID$/i</span><span class="s">)</span> &amp;&amp; <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">CompoundIDMode</span>} =~ <span class="q">/^DataField$/i</span><span class="s">)</span><span class="s">)</span> ? <span class="n">1</span> <span class="co">:</span> <span class="n">0</span><span class="sc">;</span>
+ 405   <span class="i">$CollectDataFields</span> = <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">TextOutput</span>} &amp;&amp; <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">DataFieldsMode</span>} =~ <span class="q">/^(All|Common)$/i</span><span class="s">)</span><span class="s">)</span> ? <span class="n">1</span> <span class="co">:</span> <span class="n">0</span><span class="sc">;</span>
+ 406 
+ 407   <span class="j">FILELIST:</span> <span class="k">for</span> <span class="i">$Index</span> <span class="s">(</span><span class="n">0</span> .. <span class="i">$#SDFilesList</span><span class="s">)</span> <span class="s">{</span>
+ 408     <span class="i">$SDFile</span> = <span class="i">$SDFilesList</span>[<span class="i">$Index</span>]<span class="sc">;</span>
+ 409 
+ 410     <span class="i">$SDFilesInfo</span>{<span class="w">FileOkay</span>}[<span class="i">$Index</span>] = <span class="n">0</span><span class="sc">;</span>
+ 411     <span class="i">$SDFilesInfo</span>{<span class="w">OutFileRoot</span>}[<span class="i">$Index</span>] = <span class="q">&#39;&#39;</span><span class="sc">;</span>
+ 412     <span class="i">$SDFilesInfo</span>{<span class="w">SDOutFileNames</span>}[<span class="i">$Index</span>] = <span class="q">&#39;&#39;</span><span class="sc">;</span>
+ 413     <span class="i">$SDFilesInfo</span>{<span class="w">FPOutFileNames</span>}[<span class="i">$Index</span>] = <span class="q">&#39;&#39;</span><span class="sc">;</span>
+ 414     <span class="i">$SDFilesInfo</span>{<span class="w">TextOutFileNames</span>}[<span class="i">$Index</span>] = <span class="q">&#39;&#39;</span><span class="sc">;</span>
+ 415 
+ 416     <span class="i">$SDFile</span> = <span class="i">$SDFilesList</span>[<span class="i">$Index</span>]<span class="sc">;</span>
+ 417     <span class="k">if</span> <span class="s">(</span>!<span class="s">(</span><span class="k">-e</span> <span class="i">$SDFile</span><span class="s">)</span><span class="s">)</span> <span class="s">{</span>
+ 418       <span class="k">warn</span> <span class="q">&quot;Warning: Ignoring file $SDFile: It doesn&#39;t exist\n&quot;</span><span class="sc">;</span>
+ 419       <span class="k">next</span> <span class="j">FILELIST</span><span class="sc">;</span>
+ 420     <span class="s">}</span>
+ 421     <span class="k">if</span> <span class="s">(</span>!<span class="i">CheckFileType</span><span class="s">(</span><span class="i">$SDFile</span><span class="cm">,</span> <span class="q">&quot;sd sdf&quot;</span><span class="s">)</span><span class="s">)</span> <span class="s">{</span>
+ 422       <span class="k">warn</span> <span class="q">&quot;Warning: Ignoring file $SDFile: It&#39;s not a SD file\n&quot;</span><span class="sc">;</span>
+ 423       <span class="k">next</span> <span class="j">FILELIST</span><span class="sc">;</span>
+ 424     <span class="s">}</span>
+ 425 
+ 426     <span class="k">if</span> <span class="s">(</span><span class="i">$CheckDataField</span><span class="s">)</span> <span class="s">{</span>
+ 427       <span class="c"># Make sure data field exists in SD file..</span>
+ 428       <span class="k">my</span><span class="s">(</span><span class="i">$CmpdString</span><span class="cm">,</span> <span class="i">$SpecifiedDataField</span><span class="cm">,</span> <span class="i">@CmpdLines</span><span class="cm">,</span> <span class="i">%DataFieldValues</span><span class="s">)</span><span class="sc">;</span>
+ 429 
+ 430       <span class="i">@CmpdLines</span> = <span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 431       <span class="k">open</span> <span class="w">SDFILE</span><span class="cm">,</span> <span class="q">&quot;$SDFile&quot;</span> <span class="k">or</span> <span class="k">die</span> <span class="q">&quot;Error: Couldn&#39;t open $SDFile: $! \n&quot;</span><span class="sc">;</span>
+ 432       <span class="i">$CmpdString</span> = <span class="i">ReadCmpdString</span><span class="s">(</span>\<span class="i">*SDFILE</span><span class="s">)</span><span class="sc">;</span>
+ 433       <span class="k">close</span> <span class="w">SDFILE</span><span class="sc">;</span>
+ 434       <span class="i">@CmpdLines</span> = <span class="k">split</span> <span class="q">&quot;\n&quot;</span><span class="cm">,</span> <span class="i">$CmpdString</span><span class="sc">;</span>
+ 435       <span class="i">%DataFieldValues</span> = <span class="i">GetCmpdDataHeaderLabelsAndValues</span><span class="s">(</span>\<span class="i">@CmpdLines</span><span class="s">)</span><span class="sc">;</span>
+ 436       <span class="i">$SpecifiedDataField</span> = <span class="i">$OptionsInfo</span>{<span class="w">CompoundID</span>}<span class="sc">;</span>
+ 437       <span class="k">if</span> <span class="s">(</span>!<span class="k">exists</span> <span class="i">$DataFieldValues</span>{<span class="i">$SpecifiedDataField</span>}<span class="s">)</span> <span class="s">{</span>
+ 438         <span class="k">warn</span> <span class="q">&quot;Warning: Ignoring file $SDFile: Data field value, $SpecifiedDataField, using  \&quot;--CompoundID\&quot; option in \&quot;DataField\&quot; \&quot;--CompoundIDMode\&quot; doesn&#39;t exist\n&quot;</span><span class="sc">;</span>
+ 439         <span class="k">next</span> <span class="j">FILELIST</span><span class="sc">;</span>
+ 440       <span class="s">}</span>
+ 441     <span class="s">}</span>
+ 442 
+ 443     <span class="i">$AllDataFieldsRef</span> = <span class="q">&#39;&#39;</span><span class="sc">;</span>
+ 444     <span class="i">$CommonDataFieldsRef</span> = <span class="q">&#39;&#39;</span><span class="sc">;</span>
+ 445     <span class="k">if</span> <span class="s">(</span><span class="i">$CollectDataFields</span><span class="s">)</span> <span class="s">{</span>
+ 446       <span class="k">my</span><span class="s">(</span><span class="i">$CmpdCount</span><span class="s">)</span><span class="sc">;</span>
+ 447       <span class="k">open</span> <span class="w">SDFILE</span><span class="cm">,</span> <span class="q">&quot;$SDFile&quot;</span> <span class="k">or</span> <span class="k">die</span> <span class="q">&quot;Error: Couldn&#39;t open $SDFile: $! \n&quot;</span><span class="sc">;</span>
+ 448       <span class="s">(</span><span class="i">$CmpdCount</span><span class="cm">,</span> <span class="i">$AllDataFieldsRef</span><span class="cm">,</span> <span class="i">$CommonDataFieldsRef</span><span class="s">)</span> = <span class="i">GetAllAndCommonCmpdDataHeaderLabels</span><span class="s">(</span>\<span class="i">*SDFILE</span><span class="s">)</span><span class="sc">;</span>
+ 449       <span class="k">close</span> <span class="w">SDFILE</span><span class="sc">;</span>
+ 450     <span class="s">}</span>
+ 451 
+ 452     <span class="c"># Setup output file names...</span>
+ 453     <span class="i">$FileDir</span> = <span class="q">&quot;&quot;</span><span class="sc">;</span> <span class="i">$FileName</span> = <span class="q">&quot;&quot;</span><span class="sc">;</span> <span class="i">$FileExt</span> = <span class="q">&quot;&quot;</span><span class="sc">;</span>
+ 454     <span class="s">(</span><span class="i">$FileDir</span><span class="cm">,</span> <span class="i">$FileName</span><span class="cm">,</span> <span class="i">$FileExt</span><span class="s">)</span> = <span class="i">ParseFileName</span><span class="s">(</span><span class="i">$SDFile</span><span class="s">)</span><span class="sc">;</span>
+ 455 
+ 456     <span class="i">$TextOutFileExt</span> = <span class="q">&quot;csv&quot;</span><span class="sc">;</span>
+ 457     <span class="k">if</span> <span class="s">(</span><span class="i">$Options</span>{<span class="w">outdelim</span>} =~ <span class="q">/^tab$/i</span><span class="s">)</span> <span class="s">{</span>
+ 458       <span class="i">$TextOutFileExt</span> = <span class="q">&quot;tsv&quot;</span><span class="sc">;</span>
+ 459     <span class="s">}</span>
+ 460     <span class="i">$SDOutFileExt</span> = <span class="i">$FileExt</span><span class="sc">;</span>
+ 461     <span class="i">$FPOutFileExt</span> = <span class="q">&quot;fpf&quot;</span><span class="sc">;</span>
+ 462 
+ 463     <span class="k">if</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">OutFileRoot</span>} &amp;&amp; <span class="s">(</span><span class="i">@SDFilesList</span> == <span class="n">1</span><span class="s">)</span><span class="s">)</span> <span class="s">{</span>
+ 464       <span class="k">my</span> <span class="s">(</span><span class="i">$RootFileDir</span><span class="cm">,</span> <span class="i">$RootFileName</span><span class="cm">,</span> <span class="i">$RootFileExt</span><span class="s">)</span> = <span class="i">ParseFileName</span><span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">OutFileRoot</span>}<span class="s">)</span><span class="sc">;</span>
+ 465       <span class="k">if</span> <span class="s">(</span><span class="i">$RootFileName</span> &amp;&amp; <span class="i">$RootFileExt</span><span class="s">)</span> <span class="s">{</span>
+ 466         <span class="i">$FileName</span> = <span class="i">$RootFileName</span><span class="sc">;</span>
+ 467       <span class="s">}</span>
+ 468       <span class="k">else</span> <span class="s">{</span>
+ 469         <span class="i">$FileName</span> = <span class="i">$OptionsInfo</span>{<span class="w">OutFileRoot</span>}<span class="sc">;</span>
+ 470       <span class="s">}</span>
+ 471       <span class="i">$OutFileRoot</span> = <span class="i">$FileName</span><span class="sc">;</span>
+ 472     <span class="s">}</span>
+ 473     <span class="k">else</span> <span class="s">{</span>
+ 474       <span class="i">$OutFileRoot</span> = <span class="q">&quot;${FileName}MACCSKeysFP&quot;</span><span class="sc">;</span>
+ 475     <span class="s">}</span>
+ 476 
+ 477     <span class="i">$NewSDFileName</span> = <span class="q">&quot;${OutFileRoot}.${SDOutFileExt}&quot;</span><span class="sc">;</span>
+ 478     <span class="i">$NewFPFileName</span> = <span class="q">&quot;${OutFileRoot}.${FPOutFileExt}&quot;</span><span class="sc">;</span>
+ 479     <span class="i">$NewTextFileName</span> = <span class="q">&quot;${OutFileRoot}.${TextOutFileExt}&quot;</span><span class="sc">;</span>
+ 480 
+ 481     <span class="k">if</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">SDOutput</span>}<span class="s">)</span> <span class="s">{</span>
+ 482       <span class="k">if</span> <span class="s">(</span><span class="i">$SDFile</span> =~ <span class="q">/$NewSDFileName/i</span><span class="s">)</span> <span class="s">{</span>
+ 483         <span class="k">warn</span> <span class="q">&quot;Warning: Ignoring input file $SDFile: Same output, $NewSDFileName, and input file names.\n&quot;</span><span class="sc">;</span>
+ 484         <span class="k">print</span> <span class="q">&quot;Specify a different name using \&quot;-r --root\&quot; option or use default name.\n&quot;</span><span class="sc">;</span>
+ 485         <span class="k">next</span> <span class="j">FILELIST</span><span class="sc">;</span>
+ 486       <span class="s">}</span>
+ 487     <span class="s">}</span>
+ 488 
+ 489     <span class="k">if</span> <span class="s">(</span>!<span class="i">$OptionsInfo</span>{<span class="w">OverwriteFiles</span>}<span class="s">)</span> <span class="s">{</span>
+ 490       <span class="c"># Check SD and text outout files...</span>
+ 491       <span class="k">if</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">SDOutput</span>}<span class="s">)</span> <span class="s">{</span>
+ 492         <span class="k">if</span> <span class="s">(</span><span class="k">-e</span> <span class="i">$NewSDFileName</span><span class="s">)</span> <span class="s">{</span>
+ 493           <span class="k">warn</span> <span class="q">&quot;Warning: Ignoring file $SDFile: The file $NewSDFileName already exists\n&quot;</span><span class="sc">;</span>
+ 494           <span class="k">next</span> <span class="j">FILELIST</span><span class="sc">;</span>
+ 495         <span class="s">}</span>
+ 496       <span class="s">}</span>
+ 497       <span class="k">if</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">FPOutput</span>}<span class="s">)</span> <span class="s">{</span>
+ 498         <span class="k">if</span> <span class="s">(</span><span class="k">-e</span> <span class="i">$NewFPFileName</span><span class="s">)</span> <span class="s">{</span>
+ 499           <span class="k">warn</span> <span class="q">&quot;Warning: Ignoring file $SDFile: The file $NewFPFileName already exists\n&quot;</span><span class="sc">;</span>
+ 500           <span class="k">next</span> <span class="j">FILELIST</span><span class="sc">;</span>
+ 501         <span class="s">}</span>
+ 502       <span class="s">}</span>
+ 503       <span class="k">if</span> <span class="s">(</span><span class="i">$OptionsInfo</span>{<span class="w">TextOutput</span>}<span class="s">)</span> <span class="s">{</span>
+ 504         <span class="k">if</span> <span class="s">(</span><span class="k">-e</span> <span class="i">$NewTextFileName</span><span class="s">)</span> <span class="s">{</span>
+ 505           <span class="k">warn</span> <span class="q">&quot;Warning: Ignoring file $SDFile: The file $NewTextFileName already exists\n&quot;</span><span class="sc">;</span>
+ 506           <span class="k">next</span> <span class="j">FILELIST</span><span class="sc">;</span>
+ 507         <span class="s">}</span>
+ 508       <span class="s">}</span>
+ 509     <span class="s">}</span>
+ 510 
+ 511     <span class="i">$SDFilesInfo</span>{<span class="w">FileOkay</span>}[<span class="i">$Index</span>] = <span class="n">1</span><span class="sc">;</span>
+ 512 
+ 513     <span class="i">$SDFilesInfo</span>{<span class="w">OutFileRoot</span>}[<span class="i">$Index</span>] = <span class="i">$OutFileRoot</span><span class="sc">;</span>
+ 514     <span class="i">$SDFilesInfo</span>{<span class="w">SDOutFileNames</span>}[<span class="i">$Index</span>] = <span class="i">$NewSDFileName</span><span class="sc">;</span>
+ 515     <span class="i">$SDFilesInfo</span>{<span class="w">FPOutFileNames</span>}[<span class="i">$Index</span>] = <span class="i">$NewFPFileName</span><span class="sc">;</span>
+ 516     <span class="i">$SDFilesInfo</span>{<span class="w">TextOutFileNames</span>}[<span class="i">$Index</span>] = <span class="i">$NewTextFileName</span><span class="sc">;</span>
+ 517 
+ 518     <span class="i">$SDFilesInfo</span>{<span class="w">AllDataFieldsRef</span>}[<span class="i">$Index</span>] = <span class="i">$AllDataFieldsRef</span><span class="sc">;</span>
+ 519     <span class="i">$SDFilesInfo</span>{<span class="w">CommonDataFieldsRef</span>}[<span class="i">$Index</span>] = <span class="i">$CommonDataFieldsRef</span><span class="sc">;</span>
+ 520   <span class="s">}</span>
+ 521 <span class="s">}</span>
+ 522 
+ 523 <span class="c"># Process option values...</span>
+<a name="ProcessOptions-"></a> 524 <span class="k">sub </span><span class="m">ProcessOptions</span> <span class="s">{</span>
+ 525   <span class="i">%OptionsInfo</span> = <span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 526 
+ 527   <span class="i">$OptionsInfo</span>{<span class="w">Mode</span>} = <span class="i">$Options</span>{<span class="w">mode</span>}<span class="sc">;</span>
+ 528   <span class="i">$OptionsInfo</span>{<span class="w">AromaticityModel</span>} = <span class="i">$Options</span>{<span class="w">aromaticitymodel</span>}<span class="sc">;</span>
+ 529 
+ 530   <span class="i">$OptionsInfo</span>{<span class="w">BitsOrder</span>} = <span class="i">$Options</span>{<span class="w">bitsorder</span>}<span class="sc">;</span>
+ 531   <span class="i">$OptionsInfo</span>{<span class="w">BitStringFormat</span>} = <span class="i">$Options</span>{<span class="w">bitstringformat</span>}<span class="sc">;</span>
+ 532 
+ 533   <span class="i">$OptionsInfo</span>{<span class="w">CompoundIDMode</span>} = <span class="i">$Options</span>{<span class="w">compoundidmode</span>}<span class="sc">;</span>
+ 534   <span class="i">$OptionsInfo</span>{<span class="w">CompoundIDLabel</span>} = <span class="i">$Options</span>{<span class="w">compoundidlabel</span>}<span class="sc">;</span>
+ 535   <span class="i">$OptionsInfo</span>{<span class="w">DataFieldsMode</span>} = <span class="i">$Options</span>{<span class="w">datafieldsmode</span>}<span class="sc">;</span>
+ 536 
+ 537   <span class="i">$OptionsInfo</span>{<span class="w">Filter</span>} = <span class="s">(</span><span class="i">$Options</span>{<span class="w">filter</span>} =~ <span class="q">/^Yes$/i</span><span class="s">)</span> ? <span class="n">1</span> <span class="co">:</span> <span class="n">0</span><span class="sc">;</span>
+ 538 
+ 539   <span class="k">my</span><span class="s">(</span><span class="i">@SpecifiedDataFields</span><span class="s">)</span><span class="sc">;</span>
+ 540   <span class="i">@SpecifiedDataFields</span> = <span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 541 
+ 542   <span class="i">@</span>{<span class="i">$OptionsInfo</span>{<span class="w">SpecifiedDataFields</span>}} = <span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 543   <span class="i">$OptionsInfo</span>{<span class="w">CompoundID</span>} = <span class="q">&#39;&#39;</span><span class="sc">;</span>
+ 544 
+ 545   <span class="k">if</span> <span class="s">(</span><span class="i">$Options</span>{<span class="w">datafieldsmode</span>} =~ <span class="q">/^CompoundID$/i</span><span class="s">)</span> <span class="s">{</span>
+ 546     <span class="k">if</span> <span class="s">(</span><span class="i">$Options</span>{<span class="w">compoundidmode</span>} =~ <span class="q">/^DataField$/i</span><span class="s">)</span> <span class="s">{</span>
+ 547       <span class="k">if</span> <span class="s">(</span>!<span class="i">$Options</span>{<span class="w">compoundid</span>}<span class="s">)</span> <span class="s">{</span>
+ 548         <span class="k">die</span> <span class="q">&quot;Error: You must specify a value for \&quot;--CompoundID\&quot; option in \&quot;DataField\&quot; \&quot;--CompoundIDMode\&quot;. \n&quot;</span><span class="sc">;</span>
+ 549       <span class="s">}</span>
+ 550       <span class="i">$OptionsInfo</span>{<span class="w">CompoundID</span>} = <span class="i">$Options</span>{<span class="w">compoundid</span>}<span class="sc">;</span>
+ 551     <span class="s">}</span>
+ 552     <span class="k">elsif</span> <span class="s">(</span><span class="i">$Options</span>{<span class="w">compoundidmode</span>} =~ <span class="q">/^(LabelPrefix|MolNameOrLabelPrefix)$/i</span><span class="s">)</span> <span class="s">{</span>
+ 553       <span class="i">$OptionsInfo</span>{<span class="w">CompoundID</span>} = <span class="i">$Options</span>{<span class="w">compoundid</span>} ? <span class="i">$Options</span>{<span class="w">compoundid</span>} <span class="co">:</span> <span class="q">&#39;Cmpd&#39;</span><span class="sc">;</span>
+ 554     <span class="s">}</span>
+ 555   <span class="s">}</span>
+ 556   <span class="k">elsif</span> <span class="s">(</span><span class="i">$Options</span>{<span class="w">datafieldsmode</span>} =~ <span class="q">/^Specify$/i</span><span class="s">)</span> <span class="s">{</span>
+ 557     <span class="k">if</span> <span class="s">(</span>!<span class="i">$Options</span>{<span class="w">datafields</span>}<span class="s">)</span> <span class="s">{</span>
+ 558       <span class="k">die</span> <span class="q">&quot;Error: You must specify a value for \&quot;--DataFields\&quot; option in \&quot;Specify\&quot; \&quot;-d, --DataFieldsMode\&quot;. \n&quot;</span><span class="sc">;</span>
+ 559     <span class="s">}</span>
+ 560     <span class="i">@SpecifiedDataFields</span> = <span class="k">split</span> <span class="q">/\,/</span><span class="cm">,</span> <span class="i">$Options</span>{<span class="w">datafields</span>}<span class="sc">;</span>
+ 561     <span class="k">push</span> <span class="i">@</span>{<span class="i">$OptionsInfo</span>{<span class="w">SpecifiedDataFields</span>}}<span class="cm">,</span> <span class="i">@SpecifiedDataFields</span><span class="sc">;</span>
+ 562   <span class="s">}</span>
+ 563 
+ 564   <span class="i">$OptionsInfo</span>{<span class="w">FingerprintsLabel</span>} = <span class="i">$Options</span>{<span class="w">fingerprintslabel</span>} ? <span class="i">$Options</span>{<span class="w">fingerprintslabel</span>} <span class="co">:</span> <span class="q">&#39;MACCSKeysFingerprints&#39;</span><span class="sc">;</span>
+ 565 
+ 566   <span class="i">$OptionsInfo</span>{<span class="w">KeepLargestComponent</span>} = <span class="s">(</span><span class="i">$Options</span>{<span class="w">keeplargestcomponent</span>} =~ <span class="q">/^Yes$/i</span><span class="s">)</span> ? <span class="n">1</span> <span class="co">:</span> <span class="n">0</span><span class="sc">;</span>
+ 567 
+ 568   <span class="i">$OptionsInfo</span>{<span class="w">Output</span>} = <span class="i">$Options</span>{<span class="w">output</span>}<span class="sc">;</span>
+ 569   <span class="i">$OptionsInfo</span>{<span class="w">SDOutput</span>} = <span class="s">(</span><span class="i">$Options</span>{<span class="w">output</span>} =~ <span class="q">/^(SD|All)$/i</span><span class="s">)</span> ? <span class="n">1</span> <span class="co">:</span> <span class="n">0</span><span class="sc">;</span>
+ 570   <span class="i">$OptionsInfo</span>{<span class="w">FPOutput</span>} = <span class="s">(</span><span class="i">$Options</span>{<span class="w">output</span>} =~ <span class="q">/^(FP|All)$/i</span><span class="s">)</span> ? <span class="n">1</span> <span class="co">:</span> <span class="n">0</span><span class="sc">;</span>
+ 571   <span class="i">$OptionsInfo</span>{<span class="w">TextOutput</span>} = <span class="s">(</span><span class="i">$Options</span>{<span class="w">output</span>} =~ <span class="q">/^(Text|All)$/i</span><span class="s">)</span> ? <span class="n">1</span> <span class="co">:</span> <span class="n">0</span><span class="sc">;</span>
+ 572 
+ 573   <span class="i">$OptionsInfo</span>{<span class="w">OutDelim</span>} = <span class="i">$Options</span>{<span class="w">outdelim</span>}<span class="sc">;</span>
+ 574   <span class="i">$OptionsInfo</span>{<span class="w">OutQuote</span>} = <span class="s">(</span><span class="i">$Options</span>{<span class="w">quote</span>} =~ <span class="q">/^Yes$/i</span><span class="s">)</span> ? <span class="n">1</span> <span class="co">:</span> <span class="n">0</span><span class="sc">;</span>
+ 575 
+ 576   <span class="i">$OptionsInfo</span>{<span class="w">OverwriteFiles</span>} = <span class="i">$Options</span>{<span class="w">overwrite</span>} ? <span class="n">1</span> <span class="co">:</span> <span class="n">0</span><span class="sc">;</span>
+ 577   <span class="i">$OptionsInfo</span>{<span class="w">OutFileRoot</span>} = <span class="i">$Options</span>{<span class="w">root</span>} ? <span class="i">$Options</span>{<span class="w">root</span>} <span class="co">:</span> <span class="n">0</span><span class="sc">;</span>
+ 578 
+ 579   <span class="i">$OptionsInfo</span>{<span class="w">Size</span>} = <span class="i">$Options</span>{<span class="w">size</span>}<span class="sc">;</span>
+ 580 
+ 581   <span class="i">$OptionsInfo</span>{<span class="w">VectorStringFormat</span>} = <span class="i">$Options</span>{<span class="w">vectorstringformat</span>}<span class="sc">;</span>
+ 582 <span class="s">}</span>
+ 583 
+ 584 <span class="c"># Setup script usage  and retrieve command line arguments specified using various options...</span>
+<a name="SetupScriptUsage-"></a> 585 <span class="k">sub </span><span class="m">SetupScriptUsage</span> <span class="s">{</span>
+ 586 
+ 587   <span class="c"># Retrieve all the options...</span>
+ 588   <span class="i">%Options</span> = <span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 589 
+ 590   <span class="i">$Options</span>{<span class="w">aromaticitymodel</span>} = <span class="q">&#39;MayaChemToolsAromaticityModel&#39;</span><span class="sc">;</span>
+ 591 
+ 592   <span class="i">$Options</span>{<span class="w">bitsorder</span>} = <span class="q">&#39;Ascending&#39;</span><span class="sc">;</span>
+ 593   <span class="i">$Options</span>{<span class="w">bitstringformat</span>} = <span class="q">&#39;BinaryString&#39;</span><span class="sc">;</span>
+ 594 
+ 595   <span class="i">$Options</span>{<span class="w">compoundidmode</span>} = <span class="q">&#39;LabelPrefix&#39;</span><span class="sc">;</span>
+ 596   <span class="i">$Options</span>{<span class="w">compoundidlabel</span>} = <span class="q">&#39;CompoundID&#39;</span><span class="sc">;</span>
+ 597   <span class="i">$Options</span>{<span class="w">datafieldsmode</span>} = <span class="q">&#39;CompoundID&#39;</span><span class="sc">;</span>
+ 598 
+ 599   <span class="i">$Options</span>{<span class="w">filter</span>} = <span class="q">&#39;Yes&#39;</span><span class="sc">;</span>
+ 600 
+ 601   <span class="i">$Options</span>{<span class="w">keeplargestcomponent</span>} = <span class="q">&#39;Yes&#39;</span><span class="sc">;</span>
+ 602 
+ 603   <span class="i">$Options</span>{<span class="w">mode</span>} = <span class="q">&#39;MACCSKeyBits&#39;</span><span class="sc">;</span>
+ 604 
+ 605   <span class="i">$Options</span>{<span class="w">output</span>} = <span class="q">&#39;text&#39;</span><span class="sc">;</span>
+ 606   <span class="i">$Options</span>{<span class="w">outdelim</span>} = <span class="q">&#39;comma&#39;</span><span class="sc">;</span>
+ 607   <span class="i">$Options</span>{<span class="w">quote</span>} = <span class="q">&#39;yes&#39;</span><span class="sc">;</span>
+ 608 
+ 609   <span class="i">$Options</span>{<span class="w">size</span>} = <span class="n">166</span><span class="sc">;</span>
+ 610 
+ 611   <span class="i">$Options</span>{<span class="w">vectorstringformat</span>} = <span class="q">&#39;ValuesString&#39;</span><span class="sc">;</span>
+ 612 
+ 613   <span class="k">if</span> <span class="s">(</span>!<span class="i">GetOptions</span><span class="s">(</span>\<span class="i">%Options</span><span class="cm">,</span> <span class="q">&quot;aromaticitymodel=s&quot;</span><span class="cm">,</span> <span class="q">&quot;bitsorder=s&quot;</span><span class="cm">,</span> <span class="q">&quot;bitstringformat|b=s&quot;</span><span class="cm">,</span> <span class="q">&quot;compoundid=s&quot;</span><span class="cm">,</span> <span class="q">&quot;compoundidlabel=s&quot;</span><span class="cm">,</span> <span class="q">&quot;compoundidmode=s&quot;</span><span class="cm">,</span> <span class="q">&quot;datafields=s&quot;</span><span class="cm">,</span> <span class="q">&quot;datafieldsmode|d=s&quot;</span><span class="cm">,</span> <span class="q">&quot;filter|f=s&quot;</span><span class="cm">,</span> <span class="q">&quot;fingerprintslabel=s&quot;</span><span class="cm">,</span>  <span class="q">&quot;help|h&quot;</span><span class="cm">,</span> <span class="q">&quot;keeplargestcomponent|k=s&quot;</span><span class="cm">,</span> <span class="q">&quot;mode|m=s&quot;</span><span class="cm">,</span> <span class="q">&quot;outdelim=s&quot;</span><span class="cm">,</span> <span class="q">&quot;output=s&quot;</span><span class="cm">,</span> <span class="q">&quot;overwrite|o&quot;</span><span class="cm">,</span> <span class="q">&quot;quote|q=s&quot;</span><span class="cm">,</span> <span class="q">&quot;root|r=s&quot;</span><span class="cm">,</span> <span class="q">&quot;size|s=i&quot;</span><span class="cm">,</span> <span class="q">&quot;vectorstringformat|v=s&quot;</span><span class="cm">,</span> <span class="q">&quot;workingdir|w=s&quot;</span><span class="s">)</span><span class="s">)</span> <span class="s">{</span>
+ 614     <span class="k">die</span> <span class="q">&quot;\nTo get a list of valid options and their values, use \&quot;$ScriptName -h\&quot; or\n\&quot;perl -S $ScriptName -h\&quot; command and try again...\n&quot;</span><span class="sc">;</span>
+ 615   <span class="s">}</span>
+ 616   <span class="k">if</span> <span class="s">(</span><span class="i">$Options</span>{<span class="w">workingdir</span>}<span class="s">)</span> <span class="s">{</span>
+ 617     <span class="k">if</span> <span class="s">(</span>! <span class="k">-d</span> <span class="i">$Options</span>{<span class="w">workingdir</span>}<span class="s">)</span> <span class="s">{</span>
+ 618       <span class="k">die</span> <span class="q">&quot;Error: The value specified, $Options{workingdir}, for option \&quot;-w --workingdir\&quot; is not a directory name.\n&quot;</span><span class="sc">;</span>
+ 619     <span class="s">}</span>
+ 620     <span class="k">chdir</span> <span class="i">$Options</span>{<span class="w">workingdir</span>} <span class="k">or</span> <span class="k">die</span> <span class="q">&quot;Error: Couldn&#39;t chdir $Options{workingdir}: $! \n&quot;</span><span class="sc">;</span>
+ 621   <span class="s">}</span>
+ 622   <span class="k">if</span> <span class="s">(</span>!<span class="i">Molecule::IsSupportedAromaticityModel</span><span class="s">(</span><span class="i">$Options</span>{<span class="w">aromaticitymodel</span>}<span class="s">)</span><span class="s">)</span> <span class="s">{</span>
+ 623     <span class="k">my</span><span class="s">(</span><span class="i">@SupportedModels</span><span class="s">)</span> = <span class="i">Molecule::GetSupportedAromaticityModels</span><span class="s">(</span><span class="s">)</span><span class="sc">;</span>
+ 624     <span class="k">die</span> <span class="q">&quot;Error: The value specified, $Options{aromaticitymodel}, for option \&quot;--AromaticityModel\&quot; is not valid. Supported aromaticity models in current release of MayaChemTools: @SupportedModels\n&quot;</span><span class="sc">;</span>
+ 625   <span class="s">}</span>
+ 626   <span class="k">if</span> <span class="s">(</span><span class="i">$Options</span>{<span class="w">bitsorder</span>} !~ <span class="q">/^(Ascending|Descending)$/i</span><span class="s">)</span> <span class="s">{</span>
+ 627     <span class="k">die</span> <span class="q">&quot;Error: The value specified, $Options{bitsorder}, for option \&quot;--BitsOrder\&quot; is not valid. Allowed values: Ascending or Descending\n&quot;</span><span class="sc">;</span>
+ 628   <span class="s">}</span>
+ 629   <span class="k">if</span> <span class="s">(</span><span class="i">$Options</span>{<span class="w">bitstringformat</span>} !~ <span class="q">/^(BinaryString|HexadecimalString)$/i</span><span class="s">)</span> <span class="s">{</span>
+ 630     <span class="k">die</span> <span class="q">&quot;Error: The value specified, $Options{bitstringformat}, for option \&quot;-b, --bitstringformat\&quot; is not valid. Allowed values: BinaryString or HexadecimalString\n&quot;</span><span class="sc">;</span>
+ 631   <span class="s">}</span>
+ 632   <span class="k">if</span> <span class="s">(</span><span class="i">$Options</span>{<span class="w">compoundidmode</span>} !~ <span class="q">/^(DataField|MolName|LabelPrefix|MolNameOrLabelPrefix)$/i</span><span class="s">)</span> <span class="s">{</span>
+ 633     <span class="k">die</span> <span class="q">&quot;Error: The value specified, $Options{compoundidmode}, for option \&quot;--CompoundIDMode\&quot; is not valid. Allowed values: DataField, MolName, LabelPrefix or MolNameOrLabelPrefix\n&quot;</span><span class="sc">;</span>
+ 634   <span class="s">}</span>
+ 635   <span class="k">if</span> <span class="s">(</span><span class="i">$Options</span>{<span class="w">datafieldsmode</span>} !~ <span class="q">/^(All|Common|Specify|CompoundID)$/i</span><span class="s">)</span> <span class="s">{</span>
+ 636     <span class="k">die</span> <span class="q">&quot;Error: The value specified, $Options{datafieldsmode}, for option \&quot;-d, --DataFieldsMode\&quot; is not valid. Allowed values: All, Common, Specify or CompoundID\n&quot;</span><span class="sc">;</span>
+ 637   <span class="s">}</span>
+ 638   <span class="k">if</span> <span class="s">(</span><span class="i">$Options</span>{<span class="w">filter</span>} !~ <span class="q">/^(Yes|No)$/i</span><span class="s">)</span> <span class="s">{</span>
+ 639     <span class="k">die</span> <span class="q">&quot;Error: The value specified, $Options{filter}, for option \&quot;-f, --Filter\&quot; is not valid. Allowed values: Yes or No\n&quot;</span><span class="sc">;</span>
+ 640   <span class="s">}</span>
+ 641   <span class="k">if</span> <span class="s">(</span><span class="i">$Options</span>{<span class="w">keeplargestcomponent</span>} !~ <span class="q">/^(Yes|No)$/i</span><span class="s">)</span> <span class="s">{</span>
+ 642     <span class="k">die</span> <span class="q">&quot;Error: The value specified, $Options{keeplargestcomponent}, for option \&quot;-k, --KeepLargestComponent\&quot; is not valid. Allowed values: Yes or No\n&quot;</span><span class="sc">;</span>
+ 643   <span class="s">}</span>
+ 644   <span class="k">if</span> <span class="s">(</span><span class="i">$Options</span>{<span class="w">mode</span>} !~ <span class="q">/^(MACCSKeyBits|MACCSKeyCount)$/i</span><span class="s">)</span> <span class="s">{</span>
+ 645     <span class="k">die</span> <span class="q">&quot;Error: The value specified, $Options{mode}, for option \&quot;-m, --mode\&quot; is not valid. Allowed values: MACCSKeyBits or MACCSKeyCount\n&quot;</span><span class="sc">;</span>
+ 646   <span class="s">}</span>
+ 647   <span class="k">if</span> <span class="s">(</span><span class="i">$Options</span>{<span class="w">output</span>} !~ <span class="q">/^(SD|FP|text|all)$/i</span><span class="s">)</span> <span class="s">{</span>
+ 648     <span class="k">die</span> <span class="q">&quot;Error: The value specified, $Options{output}, for option \&quot;--output\&quot; is not valid. Allowed values: SD, FP, text, or all\n&quot;</span><span class="sc">;</span>
+ 649   <span class="s">}</span>
+ 650   <span class="k">if</span> <span class="s">(</span><span class="i">$Options</span>{<span class="w">outdelim</span>} !~ <span class="q">/^(comma|semicolon|tab)$/i</span><span class="s">)</span> <span class="s">{</span>
+ 651     <span class="k">die</span> <span class="q">&quot;Error: The value specified, $Options{outdelim}, for option \&quot;--outdelim\&quot; is not valid. Allowed values: comma, tab, or semicolon\n&quot;</span><span class="sc">;</span>
+ 652   <span class="s">}</span>
+ 653   <span class="k">if</span> <span class="s">(</span><span class="i">$Options</span>{<span class="w">quote</span>} !~ <span class="q">/^(Yes|No)$/i</span><span class="s">)</span> <span class="s">{</span>
+ 654     <span class="k">die</span> <span class="q">&quot;Error: The value specified, $Options{quote}, for option \&quot;-q --quote\&quot; is not valid. Allowed values: Yes or No\n&quot;</span><span class="sc">;</span>
+ 655   <span class="s">}</span>
+ 656   <span class="k">if</span> <span class="s">(</span><span class="i">$Options</span>{<span class="w">outdelim</span>} =~ <span class="q">/semicolon/i</span> &amp;&amp; <span class="i">$Options</span>{<span class="w">quote</span>} =~ <span class="q">/^No$/i</span><span class="s">)</span> <span class="s">{</span>
+ 657     <span class="k">die</span> <span class="q">&quot;Error: The value specified, $Options{quote}, for option \&quot;-q --quote\&quot; is not allowed with, semicolon value of \&quot;--outdelim\&quot; option: Fingerprints string use semicolon as delimiter for various data fields and must be quoted.\n&quot;</span><span class="sc">;</span>
+ 658   <span class="s">}</span>
+ 659   <span class="k">if</span> <span class="s">(</span>!<span class="s">(</span><span class="i">IsPositiveInteger</span><span class="s">(</span><span class="i">$Options</span>{<span class="w">size</span>}<span class="s">)</span> &amp;&amp; <span class="s">(</span><span class="i">$Options</span>{<span class="w">size</span>} == <span class="n">166</span> || <span class="i">$Options</span>{<span class="w">size</span>} == <span class="n">322</span><span class="s">)</span><span class="s">)</span><span class="s">)</span> <span class="s">{</span>
+ 660     <span class="k">die</span> <span class="q">&quot;Error: The value specified, $Options{size}, for option \&quot;-s, --size\&quot; is not valid. Allowed values: 166 or 322 \n&quot;</span><span class="sc">;</span>
+ 661   <span class="s">}</span>
+ 662   <span class="k">if</span> <span class="s">(</span><span class="i">$Options</span>{<span class="w">vectorstringformat</span>} !~ <span class="q">/^(ValuesString|IDsAndValuesString|IDsAndValuesPairsString|ValuesAndIDsString|ValuesAndIDsPairsString)$/i</span><span class="s">)</span> <span class="s">{</span>
+ 663     <span class="k">die</span> <span class="q">&quot;Error: The value specified, $Options{vectorstringformat}, for option \&quot;-v, --VectorStringFormat\&quot; is not valid. Allowed values: ValuesString, IDsAndValuesString, IDsAndValuesPairsString, ValuesAndIDsString or ValuesAndIDsPairsString\n&quot;</span><span class="sc">;</span>
+ 664   <span class="s">}</span>
+ 665 <span class="s">}</span>
+ 666 
+<a name="EOF-"></a></pre>
+<p>&nbsp;</p>
+<br />
+<center>
+<img src="../../../images/h2o2.png">
+</center>
+</body>
+</html>