<?xml version="1.0" encoding="UTF-8"?><!DOCTYPE article PUBLIC "-//NLM//DTD JATS (Z39.96) Journal Archiving and Interchange DTD with MathML3 v1.2 20190208//EN"  "JATS-archivearticle1-mathml3.dtd"><article article-type="research-article" dtd-version="1.2" xmlns:ali="http://www.niso.org/schemas/ali/1.0/" xmlns:xlink="http://www.w3.org/1999/xlink"><front><journal-meta><journal-id journal-id-type="nlm-ta">elife</journal-id><journal-id journal-id-type="publisher-id">eLife</journal-id><journal-title-group><journal-title>eLife</journal-title></journal-title-group><issn pub-type="epub" publication-format="electronic">2050-084X</issn><publisher><publisher-name>eLife Sciences Publications, Ltd</publisher-name></publisher></journal-meta><article-meta><article-id pub-id-type="publisher-id">70597</article-id><article-id pub-id-type="doi">10.7554/eLife.70597</article-id><article-categories><subj-group subj-group-type="display-channel"><subject>Tools and Resources</subject></subj-group><subj-group subj-group-type="heading"><subject>Microbiology and Infectious Disease</subject></subj-group></article-categories><title-group><article-title>PGFinder, a novel analysis pipeline for the consistent, reproducible, and high-resolution structural analysis of bacterial peptidoglycans</article-title></title-group><contrib-group><contrib contrib-type="author" corresp="yes" id="author-240449"><name><surname>Patel</surname><given-names>Ankur V</given-names></name><contrib-id authenticated="true" contrib-id-type="orcid">https://orcid.org/0000-0001-8161-3455</contrib-id><email>apatel19@sheffield.ac.uk</email><xref ref-type="aff" rid="aff1">1</xref><xref ref-type="other" rid="fund1"/><xref ref-type="fn" rid="con1"/><xref ref-type="fn" rid="conf1"/></contrib><contrib contrib-type="author" id="author-240450"><name><surname>Turner</surname><given-names>Robert D</given-names></name><xref ref-type="aff" rid="aff2">2</xref><xref ref-type="fn" rid="con2"/><xref ref-type="fn" rid="conf1"/></contrib><contrib contrib-type="author" id="author-240451"><name><surname>Rifflet</surname><given-names>Aline</given-names></name><xref ref-type="aff" rid="aff3">3</xref><xref ref-type="aff" rid="aff4">4</xref><xref ref-type="aff" rid="aff5">5</xref><xref ref-type="other" rid="fund3"/><xref ref-type="other" rid="fund4"/><xref ref-type="other" rid="fund5"/><xref ref-type="fn" rid="con3"/><xref ref-type="fn" rid="conf1"/></contrib><contrib contrib-type="author" id="author-240452"><name><surname>Acosta-Martin</surname><given-names>Adelina E</given-names></name><xref ref-type="aff" rid="aff6">6</xref><xref ref-type="fn" rid="con4"/><xref ref-type="fn" rid="conf2"/></contrib><contrib contrib-type="author" id="author-240453"><name><surname>Nichols</surname><given-names>Andrew</given-names></name><xref ref-type="aff" rid="aff7">7</xref><xref ref-type="fn" rid="con5"/><xref ref-type="fn" rid="conf3"/></contrib><contrib contrib-type="author" id="author-240454"><name><surname>Awad</surname><given-names>Milena M</given-names></name><xref ref-type="aff" rid="aff8">8</xref><xref ref-type="other" rid="fund6"/><xref ref-type="fn" rid="con6"/><xref ref-type="fn" rid="conf1"/></contrib><contrib contrib-type="author" id="author-53246"><name><surname>Lyras</surname><given-names>Dena</given-names></name><xref ref-type="aff" rid="aff8">8</xref><xref ref-type="aff" rid="aff9">9</xref><xref ref-type="other" rid="fund6"/><xref ref-type="fn" rid="con7"/><xref ref-type="fn" rid="conf1"/></contrib><contrib contrib-type="author" id="author-240448"><name><surname>Gomperts Boneca</surname><given-names>Ivo</given-names></name><contrib-id authenticated="true" contrib-id-type="orcid">https://orcid.org/0000-0001-8122-509X</contrib-id><xref ref-type="aff" rid="aff3">3</xref><xref ref-type="aff" rid="aff4">4</xref><xref ref-type="aff" rid="aff5">5</xref><xref ref-type="other" rid="fund3"/><xref ref-type="other" rid="fund4"/><xref ref-type="other" rid="fund5"/><xref ref-type="fn" rid="con8"/><xref ref-type="fn" rid="conf1"/></contrib><contrib contrib-type="author" id="author-240455"><name><surname>Bern</surname><given-names>Marshall</given-names></name><xref ref-type="aff" rid="aff7">7</xref><xref ref-type="fn" rid="con9"/><xref ref-type="fn" rid="conf4"/></contrib><contrib contrib-type="author" corresp="yes" id="author-125931"><name><surname>Collins</surname><given-names>Mark O</given-names></name><contrib-id authenticated="true" contrib-id-type="orcid">https://orcid.org/0000-0002-7656-4975</contrib-id><email>mark.collins@sheffield.ac.uk</email><xref ref-type="aff" rid="aff1">1</xref><xref ref-type="aff" rid="aff6">6</xref><xref ref-type="fn" rid="con10"/><xref ref-type="fn" rid="conf1"/></contrib><contrib contrib-type="author" corresp="yes" id="author-97928"><name><surname>Mesnage</surname><given-names>Stéphane</given-names></name><contrib-id authenticated="true" contrib-id-type="orcid">https://orcid.org/0000-0003-1648-4890</contrib-id><email>s.mesnage@sheffield.ac.uk</email><xref ref-type="aff" rid="aff1">1</xref><xref ref-type="other" rid="fund2"/><xref ref-type="fn" rid="con11"/><xref ref-type="fn" rid="conf1"/></contrib><aff id="aff1"><label>1</label><institution>School of Biosciences, University of Sheffield</institution><addr-line><named-content content-type="city">Sheffield</named-content></addr-line><country>United Kingdom</country></aff><aff id="aff2"><label>2</label><institution>Department of Computer Science, University of Sheffield</institution><addr-line><named-content content-type="city">Sheffield</named-content></addr-line><country>United Kingdom</country></aff><aff id="aff3"><label>3</label><institution>Institut Pasteur, Unité Biologie et Génétique de la Paroi Bactérienne</institution><addr-line><named-content content-type="city">Paris</named-content></addr-line><country>France</country></aff><aff id="aff4"><label>4</label><institution>INSERM, Équipe Avenir</institution><addr-line><named-content content-type="city">Paris</named-content></addr-line><country>France</country></aff><aff id="aff5"><label>5</label><institution>CNRS, UMR 2001 &quot;Microbiologie intégrative et moléculaire&quot;</institution><addr-line><named-content content-type="city">Paris</named-content></addr-line><country>France</country></aff><aff id="aff6"><label>6</label><institution>biOMICS Facility, Faculty of Science Mass Spectrometry Centre, University of Sheffield</institution><addr-line><named-content content-type="city">Sheffield</named-content></addr-line><country>United Kingdom</country></aff><aff id="aff7"><label>7</label><institution>Protein Metrics Inc</institution><addr-line><named-content content-type="city">Cupertino</named-content></addr-line><country>United States</country></aff><aff id="aff8"><label>8</label><institution>Infection and Immunity Program, Monash Biomedicine Discovery Institute</institution><addr-line><named-content content-type="city">Clayton</named-content></addr-line><country>Australia</country></aff><aff id="aff9"><label>9</label><institution>Department of Microbiology, Monash University</institution><addr-line><named-content content-type="city">Clayton</named-content></addr-line><country>Australia</country></aff></contrib-group><contrib-group content-type="section"><contrib contrib-type="editor"><name><surname>Blokesch</surname><given-names>Melanie</given-names></name><role>Reviewing Editor</role><aff><institution>Ecole Polytechnique Fédérale de Lausanne</institution><country>Switzerland</country></aff></contrib><contrib contrib-type="senior_editor"><name><surname>Storz</surname><given-names>Gisela</given-names></name><role>Senior Editor</role><aff><institution>National Institute of Child Health and Human Development</institution><country>United States</country></aff></contrib></contrib-group><pub-date date-type="publication" publication-format="electronic"><day>28</day><month>09</month><year>2021</year></pub-date><pub-date pub-type="collection"><year>2021</year></pub-date><volume>10</volume><elocation-id>e70597</elocation-id><history><date date-type="received" iso-8601-date="2021-05-22"><day>22</day><month>05</month><year>2021</year></date><date date-type="accepted" iso-8601-date="2021-08-08"><day>08</day><month>08</month><year>2021</year></date></history><pub-history><event><event-desc>This manuscript was published as a preprint at bioRxiv.</event-desc><date date-type="preprint" iso-8601-date="2021-06-01"><day>01</day><month>06</month><year>2021</year></date><self-uri content-type="preprint" xlink:href="https://doi.org/10.1101/2021.06.01.446515"/></event></pub-history><permissions><copyright-statement>© 2021, Patel et al</copyright-statement><copyright-year>2021</copyright-year><copyright-holder>Patel et al</copyright-holder><ali:free_to_read/><license xlink:href="http://creativecommons.org/licenses/by/4.0/"><ali:license_ref>http://creativecommons.org/licenses/by/4.0/</ali:license_ref><license-p>This article is distributed under the terms of the <ext-link ext-link-type="uri" xlink:href="http://creativecommons.org/licenses/by/4.0/">Creative Commons Attribution License</ext-link>, which permits unrestricted use and redistribution provided that the original author and source are credited.</license-p></license></permissions><self-uri content-type="pdf" xlink:href="elife-70597-v1.pdf"/><self-uri content-type="figures-pdf" xlink:href="elife-70597-figures-v1.pdf"/><abstract><p>Many software solutions are available for proteomics and glycomics studies, but none are ideal for the structural analysis of peptidoglycan (PG), the essential and major component of bacterial cell envelopes. It icomprises glycan chains and peptide stems, both containing unusual amino acids and sugars. This has forced the field to rely on manual analysis approaches, which are time-consuming, labour-intensive, and prone to error. The lack of automated tools has hampered the ability to perform high-throughput analyses and prevented the adoption of a standard methodology. Here, we describe a novel tool called PGFinder for the analysis of PG structure and demonstrate that it represents a powerful tool to quantify PG fragments and discover novel structural features. Our analysis workflow, which relies on open-access tools, is a breakthrough towards a consistent and reproducible analysis of bacterial PGs. It represents a significant advance towards peptidoglycomics as a full-fledged discipline.</p></abstract><kwd-group kwd-group-type="author-keywords"><kwd>peptigoglycan</kwd><kwd>clostridium difficile</kwd><kwd>peptidoglycan structure</kwd><kwd>software</kwd><kwd>open source</kwd><kwd>jupyter notebook</kwd></kwd-group><kwd-group kwd-group-type="research-organism"><title>Research organism</title><kwd><italic>E. coli</italic></kwd><kwd><italic>C. difficile</italic></kwd></kwd-group><funding-group><award-group id="fund1"><funding-source><institution-wrap><institution-id institution-id-type="FundRef">http://dx.doi.org/10.13039/501100000268</institution-id><institution>Biotechnology and Biological Sciences Research Council</institution></institution-wrap></funding-source><award-id>BB/M011151/1</award-id><principal-award-recipient><name><surname>Patel</surname><given-names>Ankur V</given-names></name><name><surname>Mesnage</surname><given-names>Stephane</given-names></name></principal-award-recipient></award-group><award-group id="fund2"><funding-source><institution-wrap><institution-id institution-id-type="FundRef">http://dx.doi.org/10.13039/501100007155</institution-id><institution>Medical Research Council</institution></institution-wrap></funding-source><award-id>MR/S009272/1</award-id><principal-award-recipient><name><surname>Mesnage</surname><given-names>Stéphane</given-names></name></principal-award-recipient></award-group><award-group id="fund3"><funding-source><institution-wrap><institution-id institution-id-type="FundRef">http://dx.doi.org/10.13039/501100001665</institution-id><institution>Agence Nationale de la Recherche</institution></institution-wrap></funding-source><award-id>ANR-10-LABX-62-IBEID</award-id><principal-award-recipient><name><surname>Gomperts Boneca</surname><given-names>Ivo</given-names></name><name><surname>Rifflet</surname><given-names>Aline</given-names></name></principal-award-recipient></award-group><award-group id="fund4"><funding-source><institution-wrap><institution-id institution-id-type="FundRef">http://dx.doi.org/10.13039/501100001665</institution-id><institution>Agence Nationale de la Recherche</institution></institution-wrap></funding-source><award-id>ANR-16-IFEC-0004</award-id><principal-award-recipient><name><surname>Gomperts Boneca</surname><given-names>Ivo</given-names></name><name><surname>Rifflet</surname><given-names>Aline</given-names></name></principal-award-recipient></award-group><award-group id="fund5"><funding-source><institution-wrap><institution-id institution-id-type="FundRef">http://dx.doi.org/10.13039/501100001665</institution-id><institution>Agence Nationale de la Recherche</institution></institution-wrap></funding-source><award-id>ANR-18-CE15-0018</award-id><principal-award-recipient><name><surname>Gomperts Boneca</surname><given-names>Ivo</given-names></name><name><surname>Rifflet</surname><given-names>Aline</given-names></name></principal-award-recipient></award-group><award-group id="fund6"><funding-source><institution-wrap><institution-id institution-id-type="FundRef">http://dx.doi.org/10.13039/501100000923</institution-id><institution>Australian Research Council</institution></institution-wrap></funding-source><award-id>DP210103374</award-id><principal-award-recipient><name><surname>Lyras</surname><given-names>Dena</given-names></name><name><surname>Awad</surname><given-names>Milena M</given-names></name></principal-award-recipient></award-group><funding-statement>The funders had no role in study design, data collection and interpretation, or the decision to submit the work for publication.</funding-statement></funding-group><custom-meta-group><custom-meta specific-use="meta-only"><meta-name>Author impact statement</meta-name><meta-value>PGFinder is an open-source software dedicated to the analysis of peptidoglycan mass spectrometry data that paves the way for peptidoglycomics.</meta-value></custom-meta></custom-meta-group></article-meta></front><body><sec id="s1" sec-type="intro"><title>Introduction</title><p>The characterisation of bacterial cell walls started with the development of electron microscopy techniques (<xref ref-type="bibr" rid="bib20">Mudd and Lackman, 1941</xref>), and it has ever since been the focus of countless studies. The major and essential component of the bacterial cell envelope is called peptidoglycan (PG). It confers cell shape and resistance to osmotic stress and represents an unmatched target for antibiotics (<xref ref-type="bibr" rid="bib17">Mainardi et al., 2008</xref>; <xref ref-type="bibr" rid="bib27">Vollmer et al., 2008</xref>). Some of the most widely used antibiotics to date (beta-lactams and glycopeptides) inhibit the polymerisation of PG.</p><p>PG (murein; originally known as mucopeptide) is a giant, insoluble, bag-shaped molecule, and its composition was characterised soon after its discovery (<xref ref-type="bibr" rid="bib7">Cummins and Harris, 1956</xref>; <xref ref-type="bibr" rid="bib22">Rogers and Perkins, 1959</xref>; <xref ref-type="bibr" rid="bib29">Weidel and Pelzer, 1964</xref>). It is composed of glycan chains containing alternating <italic>N</italic>-acetylglucosamine (GlcNAc) and <italic>N</italic>-acetylmuramic acid (MurNAc) residues linked by β,1–4 bonds. The lactyl group of MurNAc residues is substituted by pentapeptide stems which often has the L-Ala<sub>1</sub>-γ-D-Glu<sub>2</sub>-L-DAA<sub>3</sub>-D-Ala<sub>4</sub>-D-Ala<sub>5</sub> sequence, where DAA is a diamino acid such as <italic>meso</italic>-diaminopimelic (mDAP) acid or L-lysine (<xref ref-type="fig" rid="fig1">Figure 1a</xref>; <xref ref-type="bibr" rid="bib27">Vollmer et al., 2008</xref>). In some species, a lateral chain (with variable composition and length) can be found attached to the amino acid in position 3. Peptide stem composition and polymerisation can vary amongst bacterial species (<xref ref-type="bibr" rid="bib23">Schleifer and Kandler, 1972</xref>). Whilst PG building blocks produced in the cytoplasm are always the same, the final structure undergoes constant lysis and modification, a process referred to as ‘remodelling’. Both remodelling and alternative polymerisation modes (<xref ref-type="fig" rid="fig1">Figure 1b</xref>) lead to a considerable variation in PG structure during cell growth and division. PG structural plasticity plays a critical role for adaption to environmental conditions during host-pathogen interaction (<xref ref-type="bibr" rid="bib5">Boneca et al., 2007</xref>; <xref ref-type="bibr" rid="bib14">Juan et al., 2018</xref>) or to survive exposure to antibiotics (<xref ref-type="bibr" rid="bib17">Mainardi et al., 2008</xref>).</p><fig id="fig1" position="float"><label>Figure 1.</label><caption><title>Diversity of peptidoglycan composition and structure.</title><p>(<bold>a</bold>) Representative peptidoglycan building block made of <italic>N</italic>-acetylglucosamine (GlcNAc) and <italic>N</italic>-acetylmuramic acid (MurNAc) forming a disaccharide subunit linked to a pentapeptide stem attached to the MurNAc via a lactyl moiety. Peptide stem contains both L and D-amino acids and show a great diversity in composition. Some examples of amino acids found in peptidoglycan are shown for each residue. Modifications of the sugars are also shown. (<bold>b</bold>) Representation of crosslinking diversity, 4–3 bonds (direct or via peptide crossbridge) and 3–3 bonds are made by D,D- or L,D-transpeptidases, respectively. The enzymes catalysing 1–3 and 4–2 bonds remain unknown. Acceptors stems are shown in blue and donor stems in red. DAA: diamino acid; <italic>m</italic>-DAP: <italic>meso</italic>-diaminopimelic acid; D-Lac: D-lactate; X: cell surface polymer (e.g teichoic acid); Z: lateral chain.</p></caption><graphic mime-subtype="tiff" mimetype="image" xlink:href="elife-70597-fig1-v1.tif"/></fig><p>PG material is straightforward to purify, but the structural analysis of this molecule is challenging and remains a time-consuming and labour-intensive process. The intact molecule must be broken down into soluble fragments by enzymatic digestion with a glycosyl hydrolase (lysozyme), and individual building blocks (disaccharide peptides, also called muropeptides) are analysed to gain insight into the structure of the intact molecule. A transformative step for the characterisation of disaccharide peptides has been the use of reversed-phase HPLC (rp-HPLC) and mass spectrometry (MS) towards the end of the 1990s (<xref ref-type="bibr" rid="bib10">Garcia-Bustos et al., 1988</xref>; <xref ref-type="bibr" rid="bib11">Glauner, 1988</xref>; <xref ref-type="bibr" rid="bib12">Glauner et al., 1988</xref>; <xref ref-type="bibr" rid="bib18">Martin et al., 1987</xref>). Combining muropeptides separation by rp-HPLC and MS characterisation has hinted at a more complex structure than previously reported.</p><p>Despite tremendous advances in both rp-HPLC-MS instrumentation and software development for the automated analysis of large datasets, ‘peptidoglycomics’ is still in its infancy. The experimental strategy to analyse PG structure has barely changed over the past 30 years. Even though rp-HPLC-MS has been routinely used over the past decade, the analysis of MS data remains a black box. Except for a recent study describing the PG structure of <italic>Pseudomonas aeruginosa</italic> (<xref ref-type="bibr" rid="bib2">Anderson et al., 2020b</xref>), no information is available in the literature about the strategy used to identify muropeptides in rp-HPLC-MS datasets. This task relies on searching a subset of expected structures, but the complexity of both the search space and search process is often not described.</p><p>We previously provided the proof of concept that shotgun proteomics tools can be used for the automated and unbiased analysis of PG structure (<xref ref-type="bibr" rid="bib4">Bern et al., 2017</xref>). The analysis of <italic>Clostridioides</italic> (previously <italic>Clostridium</italic>) <italic>difficile</italic> PG led to the identification of many muropeptides never reported before. This work also demonstrated that PG analysis could be carried out with relatively high throughput, opening the possibility to analyse large numbers of samples (such as clinical or environmental isolates) using a minimal amount of material (typically microgram amounts).</p><p>Here, we describe a novel software called PGFinder for the analysis of MS data. PGFinder is a versatile and straightforward open-source software tool that allows automated identification of muropeptides based on the creation of dynamic databases. Sharing PGFinder as a Jupyter Notebook provides a robust and consistent pipeline with the potential to accelerate discovery in the field of peptidoglycomics. This workflow described here allows a comprehensive description of the analysis strategy for a consistent and reproducible PG structure analysis by users in the community.</p><p>We applied the PGFinder pipeline to analyse the muropeptides composition of <italic>Escherichia coli</italic> which has been extensively studied. We demonstrate that PGFinder can capture an unprecedented level of complexity of PG structure, highlighting the limitations of the search strategies reported so far. Finally, we provide evidence that PGFinder can be used in conjunction with freely available MS data deconvolution software, making PG analysis possible using entirely open-access tools. We propose that our approach represents a significant advance towards a consistent and reproducible analysis of PG structure, allowing peptidoglycomics to take the crucial first leap to parity with other omics disciplines.</p></sec><sec id="s2" sec-type="results"><title>Results</title><sec id="s2-1"><title>PGFinder: a dedicated script for bottom-up identification of PG fragments</title><p>No pipeline is currently available for the automated analysis of MS PG data. Therefore, we sought to replicate a shotgun proteomics approach to create an analysis pipeline dedicated to PG analysis, referred to as ‘peptidoglycomics’ (<xref ref-type="bibr" rid="bib30">Wheeler et al., 2014</xref>).</p><p>To limit misidentifications due to mass coincidences, we established a search strategy relying on an iterative process (<xref ref-type="fig" rid="fig2">Figure 2</xref>). A first search was carried out using a database made of reduced disaccharide peptides (monomers) and their theoretical monoisotopic masses. MS data were deconvoluted using the Protein Metrics Byos software to generate a list of observed monoisotopic masses alongside other parameters including retention times and signal intensity. Individual theoretical masses contained in the monomer database (<xref ref-type="fig" rid="fig2">Figure 2</xref>, database 1) were compared with observed masses in the experimental dataset. Any observed mass within 10 ppm tolerance was considered as a match and the corresponding inferred structure and theoretical mass were then added to a list of matched structures (<xref ref-type="fig" rid="fig2">Figure 2</xref>, library 1). As a second step, we used the list of matched monomers to build another database in silico (<xref ref-type="fig" rid="fig2">Figure 2</xref>, database 2), corresponding to dimers and trimers and their theoretical masses. Two types of polymerisation events are included in the original PGFinder version depending on the type of crosslink either through peptide stems or glycan chains. Individual theoretical masses from the in silico database were compared to observed masses to generate a list of matched dimers and trimers (<xref ref-type="fig" rid="fig2">Figure 2</xref>, library 2). As a third step, we combined the lists of matched monomers and multimers to generate a final library of modified muropeptides (<xref ref-type="fig" rid="fig2">Figure 2</xref>, library 3). The final library contained only modified muropeptides corresponding to matched monomers, dimers, and trimers. The modifications accounted for include the presence of anhydro groups, deacetylated sugars, amidated amino acids, and modifications resulting from <italic>N</italic>-acetylglucosaminidase or amidase activities (loss of GlcNAc and lack of peptide stems, respectively). In-source decay products (loss of GlcNAc) and Na<sup>+</sup>/K<sup>+</sup> salt adducts were also added to library 3. All three libraries corresponding to observed monomers, dimers, trimers, and their modified variants were combined to search the MS data for masses matching theoretical values within a 10 ppm mass accuracy window. This search generated results processed by PGFinder to carry out a ‘clean up step’. The intensities of in-source decay products and salt adducts were combined with that from parent ions when found within close retention time (a 0.5 min time window). The output of this final step is a matched table written to a .csv format file. It contained all the inferred structures identified within the specified mass and retention time windows with an extracted-ion chromatogram (XIC) signal intensity for quantification.</p><fig id="fig2" position="float"><label>Figure 2.</label><caption><title>Flowchart outlining the algorithm for the matching script.</title><p>The identification of muropeptides was carried out using four successive steps, indicated by different colours (orange, green, blue, and red, respectively). As a first step, observed masses in the dataset are compared to a list of theoretical masses corresponding to monomers (database 1). Matched masses within the ppm tolerance set (10 ppm for Orbitrap data) are used to build a list of inferred monomeric structures and their corresponding theoretical masses (library 1). This is then used to generate a list of theoretical multimers (dimers and trimers) and their masses (database 2). A second matching round is carried out to build a list of inferred multimers (library 2). At this stage, matched monomers and multimers are combined to generate a list of modified muropeptides (library 3). Two libraries of matched theoretical masses (monomers and dimers, trimers) and a third library (their modified counterparts) are used to search the dataset. Muropeptide structures are inferred from a match within tolerance between theoretical and observed masses. This data is then ‘cleaned up’ by combining the intensities of ions corresponding to in-source decay and salt adducts to those of parent ions. The final matched mass spectrometry data is then written to a .csv file.</p></caption><graphic mime-subtype="tiff" mimetype="image" xlink:href="elife-70597-fig2-v1.tif"/></fig></sec><sec id="s2-2"><title>Using PGFinder to investigate PG structure and identify low-abundance muropeptides</title><p>The performance of the matching script was tested using the well-characterised PG from <italic>E. coli</italic> as a proof of concept. UHPLC-MS/MS data were acquired for three independent PG samples (biological replicates; <xref ref-type="fig" rid="app1fig1">Appendix 1—figure 1</xref>). Following MS1 spectral deconvolution (a process calculating masses from observed <italic>m/z</italic> values), observed masses were matched to theoretical muropeptides masses according to the strategy described above (<xref ref-type="fig" rid="fig2">Figure 2</xref>). A first search was carried out using a minimal mass database made of 10 simple PG fragments including three glycan chains (di-, tetra-, and hexasaccharides) and seven monomers (<xref ref-type="supplementary-material" rid="table1sdata1">Table 1—source data 1</xref>). Due to the .csv format of the database, the diamino acid in position 3 could not be assigned to a symbol or Greek letter and had to be one of the 26 letters already assigned by the IUPAC-IUB Joint Commission on Biochemical Nomenclature. We used the letter J for mDAP for the initial search and replaced it by the letter m in the final table. The output of the automated search is a .csv file per dataset; all files corresponding to biological replicates were collated into one Excel file (<xref ref-type="supplementary-material" rid="table1sdata2">Table 1—source data 2</xref>). Each search output contained approximately 3000 rows of masses and corresponding parameters. Depending on the dataset analysed, 41–48% of the total ion intensity was assigned to PG structures. As anticipated, inferred structures were frequently found with multiple retention times, reflecting the existence of stereoisomers, with one species accounting for most of the intensity. In some cases, observed masses matched with more than one inferred structure. The output of the automated search was consolidated as described in <xref ref-type="supplementary-material" rid="supp1">Supplementary file 1</xref>. Retention times were assigned to individual structures based on the elution of the most abundant stereoisomer. For example, &gt;97% of the most abundant monomer GM-AEJA was eluted at an average retention time of 10.04 ± 0.04 min. Data consolidation revealed an unprecedented muropeptide composition complexity compared to recent LC-MS analyses of <italic>E. coli</italic> (<xref ref-type="bibr" rid="bib15">Kühner et al., 2014</xref>; <xref ref-type="supplementary-material" rid="table1sdata2">Table 1—source data 2</xref>). Sixty PG fragments were identified (<xref ref-type="table" rid="table1">Table 1</xref>): these included glycan chains lacking peptide stems (4.38%), monomers (63.14%), dimers (29.54%), and trimers (2.94%) (<xref ref-type="fig" rid="fig3">Figure 3</xref>). Based on the abundance of multimeric PG fragments, we report a crosslinking index of 15.69%, which is slightly lower than the value previously reported of 23.1% (<xref ref-type="bibr" rid="bib11">Glauner, 1988</xref>).</p><fig id="fig3" position="float"><label>Figure 3.</label><caption><title>Distribution of <italic>E. coli</italic> peptidoglycan fragments identified using automated search workflow.</title><p>Breakdown of peptidoglycan is shown by oligomerisation state (left) branching to specific composition (right). Branch size is proportional to percentage. Monomers, dimers, trimers, and glycan chains (left) are broken down into muropeptide composition and structure (right). Individual structures are grouped by colour according to oligomerisation state. Monomers, green; dimers, yellow; trimers, orange. Residues in square brackets are only found in some muropeptides. For example, GM-AEJ[A]-GM-AEJ[A] can represent GM-AEJA-GM-AEJA, GM-AEJA-GM-AEJ, and GM-AEJ-GM-AEJ. G: <italic>N</italic>-acetylglucosamine; M, <italic>N-</italic>acetylmuramic acid; A: L- or D- alanine; E: γ-D-glutamic acid; J: <italic>meso</italic>-diaminopimelic acid; K: D-lysine; R: D-arginine; G: glycine.</p></caption><graphic mime-subtype="tiff" mimetype="image" xlink:href="elife-70597-fig3-v1.tif"/></fig><table-wrap id="table1" position="float"><label>Table 1.</label><caption><title>Processed match output.</title><p><supplementary-material id="table1sdata1"><label>Table 1—source data 1.</label><caption><title><italic>E. coli</italic> simple mass database.</title></caption><media mime-subtype="octet-stream" mimetype="application" xlink:href="elife-70597-table1-data1-v1.csv"/></supplementary-material><supplementary-material id="table1sdata2"><label>Table 1—source data 2.</label><caption><title><italic>E. coli</italic> matching output and consolidated data.</title></caption><media mime-subtype="xlsx" mimetype="application" xlink:href="elife-70597-table1-data2-v1.xlsx"/></supplementary-material><supplementary-material id="table1sdata3"><label>Table 1—source data 3.</label><caption><title>MS/MS analysis of <italic>E. coli</italic> glycan chains and monomers.</title></caption><media mime-subtype="pdf" mimetype="application" xlink:href="elife-70597-table1-data3-v1.pdf"/></supplementary-material><supplementary-material id="table1sdata4"><label>Table 1—source data 4.</label><caption><title><italic>E. coli</italic> complex mass database.</title></caption><media mime-subtype="octet-stream" mimetype="application" xlink:href="elife-70597-table1-data4-v1.csv"/></supplementary-material><supplementary-material id="table1sdata5"><label>Table 1—source data 5.</label><caption><title><italic>E. coli</italic> muropeptide complex table.</title></caption><media mime-subtype="xlsx" mimetype="application" xlink:href="elife-70597-table1-data5-v1.xlsx"/></supplementary-material></p></caption><table frame="hsides" rules="groups"><thead><tr><th align="left" rowspan="2" valign="bottom"/><th align="left" rowspan="2" valign="bottom">Structure</th><th align="left" valign="bottom">RT (min)</th><th align="left" valign="bottom">Abundance (%)</th><th align="left" colspan="3" valign="bottom">Monoisotopic mass (Da)</th></tr><tr><th align="left" valign="bottom">Av±SD</th><th align="left" valign="bottom">Av±SD</th><th align="left" valign="bottom">Obs</th><th align="left" valign="bottom">Theo</th><th align="left" valign="bottom">Δppm</th></tr></thead><tbody><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM|0</td><td align="char" char="." valign="bottom">3.62±0.01</td><td align="char" char="." valign="bottom">3.465±0.683</td><td align="char" char="." valign="bottom">498.205</td><td align="char" char="." valign="bottom">498.206</td><td align="char" char="." valign="bottom">2.5</td></tr><tr><td align="left" valign="bottom">Glycans</td><td align="left" valign="bottom">GM (x2)|0</td><td align="char" char="." valign="bottom">10.11±0.03</td><td align="char" char="." valign="bottom">0.428±0.349</td><td align="char" char="." valign="bottom">976.384</td><td align="char" char="." valign="bottom">976.386</td><td align="char" char="." valign="bottom">2.2</td></tr><tr><td align="left" valign="bottom">4.38%±0.35%</td><td align="left" valign="bottom">GM (anhydro) |0</td><td align="char" char="." valign="bottom">8.20±1.92</td><td align="char" char="." valign="bottom">0.238±0.025</td><td align="char" char="." valign="bottom">478.179</td><td align="char" char="." valign="bottom">478.180</td><td align="char" char="." valign="bottom">2.9</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM (deacetyl) |0</td><td align="char" char="." valign="bottom">2.57±0.00</td><td align="char" char="." valign="bottom">0.155±0.032</td><td align="char" char="." valign="bottom">456.194</td><td align="char" char="." valign="bottom">456.196</td><td align="char" char="." valign="bottom">3.5</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM (x2) (deacetyl) |0</td><td align="char" char="." valign="bottom">6.86±0.02</td><td align="char" char="." valign="bottom">0.093±0.012</td><td align="char" char="." valign="bottom">934.372</td><td align="char" char="." valign="bottom">934.376</td><td align="char" char="." valign="bottom">3.2</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA|1</td><td align="char" char="." valign="bottom">10.04±0.04</td><td align="char" char="." valign="bottom">36.098±2.131</td><td align="char" char="." valign="bottom">941.405</td><td align="char" char="." valign="bottom">941.408</td><td align="char" char="." valign="bottom">2.8</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEm|1</td><td align="char" char="." valign="bottom">6.57±0.01</td><td align="char" char="." valign="bottom">14.352±0.397</td><td align="char" char="." valign="bottom">870.368</td><td align="char" char="." valign="bottom">870.371</td><td align="char" char="." valign="bottom">3.0</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmKR|1</td><td align="char" char="." valign="bottom">9.56±0.05</td><td align="char" char="." valign="bottom">8.030±0.774</td><td align="char" char="." valign="bottom">1154.563</td><td align="char" char="." valign="bottom">1154.567</td><td align="char" char="." valign="bottom">3.6</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AE|1</td><td align="char" char="." valign="bottom">9.57±0.04</td><td align="char" char="." valign="bottom">1.809±0.231</td><td align="char" char="." valign="bottom">698.284</td><td align="char" char="." valign="bottom">698.286</td><td align="char" char="." valign="bottom">3.1</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmG|1</td><td align="char" char="." valign="bottom">7.85±0.05</td><td align="char" char="." valign="bottom">0.689±0.049</td><td align="char" char="." valign="bottom">927.390</td><td align="char" char="." valign="bottom">927.392</td><td align="char" char="." valign="bottom">2.3</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEm (anhydro) |1</td><td align="char" char="." valign="bottom">13.98±0.02</td><td align="char" char="." valign="bottom">0.668±0.073</td><td align="char" char="." valign="bottom">850.342</td><td align="char" char="." valign="bottom">850.344</td><td align="char" char="." valign="bottom">2.2</td></tr><tr><td align="left" valign="bottom">Monomers</td><td align="left" valign="bottom">GM-AEmA (anhydro) |1</td><td align="char" char="." valign="bottom">16.55±0.01</td><td align="char" char="." valign="bottom">0.573±0.100</td><td align="char" char="." valign="bottom">921.380</td><td align="char" char="." valign="bottom">921.382</td><td align="char" char="." valign="bottom">2.0</td></tr><tr><td align="left" valign="bottom">63.14%±1.13%</td><td align="left" valign="bottom">GM-AEmAG|1</td><td align="char" char="." valign="bottom">9.45±0.05</td><td align="char" char="." valign="bottom">0.219±0.009</td><td align="char" char="." valign="bottom">998.426</td><td align="char" char="." valign="bottom">998.429</td><td align="char" char="." valign="bottom">3.1</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmKR (anhydro) |1</td><td align="char" char="." valign="bottom">14.83±0.01</td><td align="char" char="." valign="bottom">0.160±0.039</td><td align="char" char="." valign="bottom">1134.537</td><td align="char" char="." valign="bottom">1134.540</td><td align="char" char="." valign="bottom">2.9</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA (deacetyl) |1</td><td align="char" char="." valign="bottom">8.57±0.06</td><td align="char" char="." valign="bottom">0.083±0.055</td><td align="char" char="." valign="bottom">899.394</td><td align="char" char="." valign="bottom">899.397</td><td align="char" char="." valign="bottom">3.1</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-GM-AEmA|1</td><td align="char" char="." valign="bottom">13.10±0.02</td><td align="char" char="." valign="bottom">0.075±0.040</td><td align="char" char="." valign="bottom">1419.584</td><td align="char" char="." valign="bottom">1419.588</td><td align="char" char="." valign="bottom">2.9</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AE (anhydro) |1</td><td align="char" char="." valign="bottom">17.44±0.01</td><td align="char" char="." valign="bottom">0.069±0.013</td><td align="char" char="." valign="bottom">678.258</td><td align="char" char="." valign="bottom">678.260</td><td align="char" char="." valign="bottom">2.8</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">M-AEm|1</td><td align="char" char="." valign="bottom">4.56±0.01</td><td align="char" char="." valign="bottom">0.062±0.064</td><td align="char" char="." valign="bottom">667.289</td><td align="char" char="." valign="bottom">667.291</td><td align="char" char="." valign="bottom">3.8</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">M-AEmKR|1</td><td align="char" char="." valign="bottom">8.16±0.06</td><td align="char" char="." valign="bottom">0.061±0.056<xref ref-type="table-fn" rid="table1fn3">*</xref></td><td align="char" char="." valign="bottom">951.484</td><td align="char" char="." valign="bottom">951.487</td><td align="char" char="." valign="bottom">3.2</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmAA|1</td><td align="char" char="." valign="bottom">11.38±0.04</td><td align="char" char="." valign="bottom">0.059±0.003</td><td align="char" char="." valign="bottom">1012.442</td><td align="char" char="." valign="bottom">1012.445</td><td align="char" char="." valign="bottom">2.4</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">M-AEmA|1</td><td align="char" char="." valign="bottom">8.52±0.05</td><td align="char" char="." valign="bottom">0.053±0.015</td><td align="char" char="." valign="bottom">738.325</td><td align="char" char="." valign="bottom">738.328</td><td align="char" char="." valign="bottom">4.0</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-GM-AEm|1</td><td align="char" char="." valign="bottom">11.31±0.04</td><td align="char" char="." valign="bottom">0.042±0.025</td><td align="char" char="." valign="bottom">1348.547</td><td align="char" char="." valign="bottom">1348.551</td><td align="char" char="." valign="bottom">2.4</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEm (deacetyl) |1</td><td align="char" char="." valign="bottom">4.77±0.01</td><td align="char" char="." valign="bottom">0.024±0.014</td><td align="char" char="." valign="bottom">828.358</td><td align="char" char="." valign="bottom">828.360</td><td align="char" char="." valign="bottom">3.0</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-GM-AEmKR|1</td><td align="char" char="." valign="bottom">12.18±0.03</td><td align="char" char="." valign="bottom">0.011±0.002<xref ref-type="table-fn" rid="table1fn3">*</xref></td><td align="char" char="." valign="bottom">1632.742</td><td align="char" char="." valign="bottom">1632.747</td><td align="char" char="." valign="bottom">3.0</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA-GM-AEmA|2</td><td align="char" char="." valign="bottom">16.01±0.02</td><td align="char" char="." valign="bottom">17.247±0.777</td><td align="char" char="." valign="bottom">1864.800</td><td align="char" char="." valign="bottom">1864.805</td><td align="char" char="." valign="bottom">2.3</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA-GM-AEmKR|2</td><td align="char" char="." valign="bottom">14.83±0.02</td><td align="char" char="." valign="bottom">4.589±0.589</td><td align="char" char="." valign="bottom">2077.957</td><td align="char" char="." valign="bottom">2077.964</td><td align="char" char="." valign="bottom">3.0</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA-GM-AEm|2</td><td align="char" char="." valign="bottom">15.09±0.02</td><td align="char" char="." valign="bottom">3.207±0.168</td><td align="char" char="." valign="bottom">1793.763</td><td align="char" char="." valign="bottom">1793.768</td><td align="char" char="." valign="bottom">2.6</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA-GM-AEmA (anhydro) |2</td><td align="char" char="." valign="bottom">20.56±0.01</td><td align="char" char="." valign="bottom">0.873±0.037</td><td align="char" char="." valign="bottom">1844.774</td><td align="char" char="." valign="bottom">1844.778</td><td align="char" char="." valign="bottom">2.4</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEm-GM-AEmKR|2</td><td align="char" char="." valign="bottom">14.22±0.00</td><td align="char" char="." valign="bottom">0.855±0.101</td><td align="char" char="." valign="bottom">2006.920</td><td align="char" char="." valign="bottom">2006.926</td><td align="char" char="." valign="bottom">3.3</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA-GM-AEmKR (anhydro) |2</td><td align="char" char="." valign="bottom">18.89±0.17</td><td align="char" char="." valign="bottom">0.665±0.079</td><td align="char" char="." valign="bottom">2057.934</td><td align="char" char="." valign="bottom">2057.937</td><td align="char" char="." valign="bottom">1.8</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEm-GM-AEm|2</td><td align="char" char="." valign="bottom">14.23±0.01</td><td align="char" char="." valign="bottom">0.558±0.062</td><td align="char" char="." valign="bottom">1722.725</td><td align="char" char="." valign="bottom">1722.730</td><td align="char" char="." valign="bottom">3.0</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEm-GM-AEmAG|2</td><td align="char" char="." valign="bottom">14.68±0.01</td><td align="char" char="." valign="bottom">0.416±0.025</td><td align="char" char="." valign="bottom">1850.785</td><td align="char" char="." valign="bottom">1850.789</td><td align="char" char="." valign="bottom">2.4</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA-GM-AEm (anhydro) |2</td><td align="char" char="." valign="bottom">19.66±0.01</td><td align="char" char="." valign="bottom">0.381±0.028</td><td align="char" char="." valign="bottom">1773.738</td><td align="char" char="." valign="bottom">1773.741</td><td align="char" char="." valign="bottom">2.1</td></tr><tr><td align="left" valign="bottom">Dimers</td><td align="left" valign="bottom">GM-AEmA-GM-AEmAG|2</td><td align="char" char="." valign="bottom">15.33±0.02</td><td align="char" char="." valign="bottom">0.179±0.005</td><td align="char" char="." valign="bottom">1921.822</td><td align="char" char="." valign="bottom">1921.826</td><td align="char" char="." valign="bottom">2.2</td></tr><tr><td align="left" valign="bottom">29.54%±0.46%</td><td align="left" valign="bottom">GM-AEm-GM-AEmKR (anhydro) |2</td><td align="char" char="." valign="bottom">18.07±0.01</td><td align="char" char="." valign="bottom">0.170±0.024</td><td align="char" char="." valign="bottom">1986.896</td><td align="char" char="." valign="bottom">1986.900</td><td align="char" char="." valign="bottom">2.1</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEm-GM-AEm (anhydro) |2</td><td align="char" char="." valign="bottom">18.77±0.01</td><td align="char" char="." valign="bottom">0.141±0.015</td><td align="char" char="." valign="bottom">1702.697</td><td align="char" char="." valign="bottom">1702.704</td><td align="char" char="." valign="bottom">4.5</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA-GM-AEmAA|2</td><td align="char" char="." valign="bottom">16.54±0.01</td><td align="char" char="." valign="bottom">0.075±0.002</td><td align="char" char="." valign="bottom">1935.838</td><td align="char" char="." valign="bottom">1935.842</td><td align="char" char="." valign="bottom">2.1</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEm-GM-AEmG|2</td><td align="char" char="." valign="bottom">13.91±0.01</td><td align="char" char="." valign="bottom">0.054±0.003</td><td align="char" char="." valign="bottom">1779.747</td><td align="char" char="." valign="bottom">1779.752</td><td align="char" char="." valign="bottom">2.7</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-GM-AEmA-GM-AEmA|2</td><td align="char" char="." valign="bottom">17.51±0.01</td><td align="char" char="." valign="bottom">0.046±0.028</td><td align="char" char="." valign="bottom">2342.976</td><td align="char" char="." valign="bottom">2342.985</td><td align="char" char="." valign="bottom">3.6</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA-GM-AEmA (deacetyl) |2</td><td align="char" char="." valign="bottom">15.17±0.01</td><td align="char" char="." valign="bottom">0.029±0.022</td><td align="char" char="." valign="bottom">1822.789</td><td align="char" char="." valign="bottom">1822.794</td><td align="char" char="." valign="bottom">3.0</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA-GM-AEmG (anhydro) |2</td><td align="char" char="." valign="bottom">19.12±0.01</td><td align="char" char="." valign="bottom">0.021±0.001</td><td align="char" char="." valign="bottom">1830.761</td><td align="char" char="." valign="bottom">1830.763</td><td align="char" char="." valign="bottom">0.7</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA-GM-AEmAG (anhydro) |2</td><td align="char" char="." valign="bottom">19.73±0.01</td><td align="char" char="." valign="bottom">0.019±0.002</td><td align="char" char="." valign="bottom">1901.796</td><td align="char" char="." valign="bottom">1901.800</td><td align="char" char="." valign="bottom">2.1</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA-GM-AEmAA (anhydro) |2</td><td align="char" char="." valign="bottom">21.17±0.02</td><td align="char" char="." valign="bottom">0.015±0.002</td><td align="char" char="." valign="bottom">1915.812</td><td align="char" char="." valign="bottom">1915.816</td><td align="char" char="." valign="bottom">1.8</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-GM-AEmA-GM-AEm|2</td><td align="char" char="." valign="bottom">16.85±0.00</td><td align="char" char="." valign="bottom">0.003±0.004</td><td align="char" char="." valign="bottom">2271.943</td><td align="char" char="." valign="bottom">2271.947</td><td align="char" char="." valign="bottom">2.1</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA-GM-AEmA-GM-AEmA|3</td><td align="char" char="." valign="bottom">18.86±0.01</td><td align="char" char="." valign="bottom">1.751±0.221</td><td align="char" char="." valign="bottom">2788.192</td><td align="char" char="." valign="bottom">2788.202</td><td align="char" char="." valign="bottom">3.5</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA-GM-AEmA-GM-AEm|3</td><td align="char" char="." valign="bottom">18.23±0.21</td><td align="char" char="." valign="bottom">0.371±0.031</td><td align="char" char="." valign="bottom">2717.158</td><td align="char" char="." valign="bottom">2717.164</td><td align="char" char="." valign="bottom">2.2</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA-GM-AEmA-GM-AEmA (anhydro) |3</td><td align="char" char="." valign="bottom">22.39±0.02</td><td align="char" char="." valign="bottom">0.222±0.027</td><td align="char" char="." valign="bottom">2768.169</td><td align="char" char="." valign="bottom">2768.175</td><td align="char" char="." valign="bottom">2.3</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA-GM-AEmA-GM-AEmKR|3</td><td align="char" char="." valign="bottom">17.54±0.01</td><td align="char" char="." valign="bottom">0.207±0.028</td><td align="char" char="." valign="bottom">3001.350</td><td align="char" char="." valign="bottom">3001.360</td><td align="char" char="." valign="bottom">3.4</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA-GM-AEmA-GM-AEm (anhydro) |3</td><td align="char" char="." valign="bottom">21.60±0.02</td><td align="char" char="." valign="bottom">0.117±0.003</td><td align="char" char="." valign="bottom">2697.133</td><td align="char" char="." valign="bottom">2697.138</td><td align="char" char="." valign="bottom">1.8</td></tr><tr><td align="left" valign="bottom">Trimers</td><td align="left" valign="bottom">GM-AEmA-GM-AEmA-GM-AEmKR (anhydro) |3</td><td align="char" char="." valign="bottom">20.90±0.16</td><td align="char" char="." valign="bottom">0.088±0.026</td><td align="char" char="." valign="bottom">2981.328</td><td align="char" char="." valign="bottom">2981.334</td><td align="char" char="." valign="bottom">2.2</td></tr><tr><td align="left" valign="bottom">2.94%±0.36%</td><td align="left" valign="bottom">GM-AEmA-GM-AEmA-GM-AEmG|3</td><td align="char" char="." valign="bottom">17.72±0.01</td><td align="char" char="." valign="bottom">0.039±0.004</td><td align="char" char="." valign="bottom">2774.182</td><td align="char" char="." valign="bottom">2774.186</td><td align="char" char="." valign="bottom">1.4</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA-GM-AEm-GM-AEm|3</td><td align="char" char="." valign="bottom">17.45±0.01</td><td align="char" char="." valign="bottom">0.029±0.005</td><td align="char" char="." valign="bottom">2646.123</td><td align="char" char="." valign="bottom">2646.127</td><td align="char" char="." valign="bottom">1.7</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA-GM-AEm-GM-AEm (anhydro) |3</td><td align="char" char="." valign="bottom">21.16±0.01</td><td align="char" char="." valign="bottom">0.025±0.001</td><td align="char" char="." valign="bottom">2626.096</td><td align="char" char="." valign="bottom">2626.101</td><td align="char" char="." valign="bottom">1.9</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA-GM-AEm-GM-AEmKR|3</td><td align="char" char="." valign="bottom">17.11±0.01</td><td align="char" char="." valign="bottom">0.022±0.002</td><td align="char" char="." valign="bottom">2930.316</td><td align="char" char="." valign="bottom">2930.323</td><td align="char" char="." valign="bottom">2.7</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA-GM-AEmA-GM-AEmAG|3</td><td align="char" char="." valign="bottom">18.24±0.01</td><td align="char" char="." valign="bottom">0.021±0.001</td><td align="char" char="." valign="bottom">2845.217</td><td align="char" char="." valign="bottom">2845.223</td><td align="char" char="." valign="bottom">2.0</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA-GM-AEmA-GM-AEmAA|3</td><td align="char" char="." valign="bottom">19.23±0.01</td><td align="char" char="." valign="bottom">0.014±0.002<xref ref-type="table-fn" rid="table1fn3">*</xref></td><td align="char" char="." valign="bottom">2859.235</td><td align="char" char="." valign="bottom">2859.239</td><td align="char" char="." valign="bottom">1.3</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA-GM-AEm-GM-AEmKR (anhydro) |3</td><td align="char" char="." valign="bottom">20.31±0.02</td><td align="char" char="." valign="bottom">0.014±0.005</td><td align="char" char="." valign="bottom">2910.293</td><td align="char" char="." valign="bottom">2910.297</td><td align="char" char="." valign="bottom">1.5</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEm-GM-AEmG-GM-AEmAG|3</td><td align="char" char="." valign="bottom">17.18±0.00</td><td align="char" char="." valign="bottom">0.004±0.005</td><td align="char" char="." valign="bottom">2703.143</td><td align="char" char="." valign="bottom">2703.149</td><td align="char" char="." valign="bottom">2.0</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA-GM-AEm-GM-AEmG (anhydro) |3</td><td align="char" char="." valign="bottom">21.21±0.02</td><td align="char" char="." valign="bottom">0.011±0.003<xref ref-type="table-fn" rid="table1fn3">*</xref></td><td align="char" char="." valign="bottom">2754.157</td><td align="char" char="." valign="bottom">2754.160</td><td align="char" char="." valign="bottom">1.1</td></tr><tr><td align="left" valign="bottom"/><td align="left" valign="bottom">GM-AEmA-GM-AEmA-GM-AEmAG (anhydro) |3</td><td align="char" char="." valign="bottom">21.77±0.01</td><td align="char" char="." valign="bottom">0.006±0.004</td><td align="char" char="." valign="bottom">2825.189</td><td align="char" char="." valign="bottom">2825.197</td><td align="char" char="." valign="bottom">2.8</td></tr></tbody></table><table-wrap-foot><fn><p>Inferrred dimers and trimers are based on the most abundant monomers and could correspond to alternative structures.</p></fn><fn><p>G: GlcNAc; M: MurNAc; m: <italic>meso</italic>-diaminopimelic acid; the number following the symbol ‘|’ refers to the oligomerisation state (1 for monomers, 2 for dimers, and 3 for trimers).</p></fn><fn id="table1fn3"><label>*</label><p>Calculated from two values.</p></fn></table-wrap-foot></table-wrap><p>The automated and unbiased search revealed several muropeptides that were expected but never reported to date for <italic>E. coli</italic>. These included (i) PG fragments resulting from amidase activity (4.55%), found as ‘denuded glycans’ (disaccharides and tetrasaccharides) and modified variants or muropeptide stem with an extra disaccharide residue; (ii) a low-abundance (0.23%) PG fragments containing deacetylated GlcNAc residues; and (iii) PG fragments resulting from glucosaminidase activity (0.12%). Deacetylated muropeptides were not expected since no <italic>E. coli</italic> PG deacetylase has been identified in this organism to date. All the structures identified for the first time (glycan chains, monomers containing deacetyl groups, and muropeptides lacking a GlcNAc residue) were confirmed by MS/MS analysis (<xref ref-type="supplementary-material" rid="table1sdata3">Table 1—source data 3</xref>). The proportion of muropeptides with anhydroMurNAc groups identified (4.55%) was in line with previous studies, yielding an average chain length of 36.05. This value is higher than an earlier study that reported a predominant chain length of 5–10 disaccharide units (<xref ref-type="bibr" rid="bib13">Harz et al., 1990</xref>), but agreed with recent work that reported long glycan chains in <italic>E. coli</italic> (<xref ref-type="bibr" rid="bib25">Turner et al., 2018</xref>). Overall, the quantification of muropeptides across biological replicates was very consistent, with Pearson’s correlation coefficients &gt;0.96 (<xref ref-type="fig" rid="app1fig2">Appendix 1—figure 2</xref>). The most pronounced variations in quantification were observed with low-abundance muropeptides accounting for less than 1% of the species identified.</p><p>We further explored the structural diversity of <italic>E. coli</italic> PG, performing a more complex search with a mass database made of glycan chains and all possible monomers containing di-, tri-, tetra-, and pentapeptide stems (<xref ref-type="supplementary-material" rid="table1sdata4">Table 1—source data 4</xref>; 224 structures in total). Several monomers with tetra- and pentapeptide stems containing unusual amino acids were identified. Only four structures could be confirmed by MS/MS analysis and corresponding multimers were retained (GM-AEJF, -AEJN, AEJK, -AEJAK, and GM-AEJKD). Collectively, muropeptides containing unusual amino acids accounted for ca. 7.5% of the 80 structures identified (<xref ref-type="supplementary-material" rid="table1sdata5">Table 1—source data 5</xref>) using a complex mass database.</p></sec><sec id="s2-3"><title>Using PGFinder for the comparative analysis of PG structures</title><p>We showed with <italic>E. coli</italic> data that PGFinder is suitable to characterise the high-resolution structure of PGs using a ‘bottom-up’ approach. However, this requires a careful analysis of the search output to confirm the identity of muropeptides identified and discriminate between multiple structures that can be assigned to a unique observed mass. A more basic application is the use of PGFinder in organisms that have already been studied in detail to either compare PG composition or quantify the abundance of specific structures. This application accounts for most PG analyses described in the literature, comparing PG structures between different isolates, isogenic mutants, or cells grown in different conditions.</p><p>We chose the PG of <italic>C. difficile</italic> as a proof of concept to demonstrate how PGFinder can carry out a straightforward comparative analysis. We prepared PG samples corresponding to biological triplicates from two clinical isolates, R20291 and M7404. To illustrate the versatility of the software, PG samples were digested by mutanolysin, and disaccharide peptides were beta-eliminated to generate lactyl-peptides (<xref ref-type="bibr" rid="bib24">Tipper, 2002</xref>) and analysed by UHPLC-MS (<xref ref-type="fig" rid="fig4s1">Figure 4—figure supplement 1</xref>). Since the high-resolution PG structure of <italic>C. difficile</italic> has been described based on the MS/MS analysis of muropeptides (<xref ref-type="bibr" rid="bib4">Bern et al., 2017</xref>), we used data previously published to build a PG fragment database for a ‘one-off’ matching step. The database included monomers identified by MS/MS, containing unusual amino acids and the corresponding dimers resulting from either D,D- or L,D-transpeptidation (with a AEJA or AEJ peptide as donor stem). To limit the complexity of the search output, we limited the list of trimers, tetramers, and pentamers to those containing the most abundant peptide stems found in dimers (AEJA, AEJ, or AEJG). The database of theoretical masses contained 74 PG structures described in <xref ref-type="supplementary-material" rid="fig4sdata1">Figure 4—source data 1</xref>. MS data were deconvoluted using Byos Feature Finder, and observed monoisotopic masses were matched to theoretical masses using PGFinder. To perform the matching operation in its simplest form, all options offered by the software were deactivated.</p><p>All monomers and dimers searched were identified, except for two muropeptides containing a peptide stem with a methionine residue in position four, both present in low abundance in <italic>C. difficile</italic> strain 630 (ca. 0.10%). 25 out of the 31 possible trimers, tetramers, and pentamers searched were found (<xref ref-type="supplementary-material" rid="fig4sdata2">Figure 4—source data 2</xref>). Comparison of biological replicates revealed a high reproducibility, both in retention times and quantification. A high correlation was found between biological replicates, confirming the robustness of the quantification method (<xref ref-type="fig" rid="fig4">Figure 4a</xref>). Both strains contained a similar amount of mono-, tri-, tetra-, and pentamers, but strain M7404 contained a significantly lower proportion of monomers and a higher proportion of dimers than strain R20291 (p=0.022 and p=0.005, respectively; Student’s <italic>t</italic>-test; <xref ref-type="fig" rid="fig4">Figure 4b</xref>). We next performed a Student’s <italic>t</italic>-test using permutation-based FDR to identify statistically significant differences in the abundance of individual muropeptides between the two strains. The p-value was plotted on a volcano plot against the fold change in abundance between the two samples (<xref ref-type="fig" rid="fig4">Figure 4c</xref>). Two muropeptides were significantly less abundant (Lac-AEJ[AG] and Lac-AEJ-Lac-AEJA), and four others were significantly more abundant in strain R20291 (Lac-AEJV- Lac-AEJA, Lac-AEJ[L/I]-Lac-AEJA, Lac-AEJAA-Lac-AEJA and the trimer (Lac-AEJA)3). These differences are likely to reflect different substrate specificities for PBPs and the Ddl ligases in these strains. Therefore, combining the output of PGFinder with statistical analysis of muropeptide abundance offers a robust workflow to identify differences in PG composition.</p><fig-group><fig id="fig4" position="float"><label>Figure 4.</label><caption><title>Comparative analysis of <italic>C. difficile</italic> R20291 and M7404 peptidoglycan (PG) composition.</title><p>(<bold>a</bold>) Pearson’s correlation coefficients across biological replicates of R20291 and M7404 <italic>C. difficile</italic> isolates. Heatmap gradient shows highest value in green to lowest value in red. (<bold>b</bold>) Muropeptide distribution according to degree of crosslinking. Comparison was carried out using a Student’s <italic>t</italic>-test; p-value is indicated for each category of muropeptides. (<bold>c</bold>) Volcano plot, where each dot represents an individual muropeptide, plotted against the significance (Student’s <italic>t</italic>-test p-value&lt;0.05, FDR &lt; 0.05, S<sub>0</sub> = 0.1) and difference (log<sub>2</sub>). Muropeptides showing a significantly different abundance between strains are highlighted in red. Lac: lactyl group; A: D/L-alanine; E: γ-D-glutamate; J: <italic>meso</italic>-diaminopimelic acid V: D-valine; L: D-leucine; I: D-isoleucine; G: glycine.</p><p><supplementary-material id="fig4sdata1"><label>Figure 4—source data 1.</label><caption><title><italic>C.</italic> <italic>difficile</italic> mass database.</title></caption><media mime-subtype="csv" mimetype="application" xlink:href="elife-70597-fig4-data1-v1.csv"/></supplementary-material></p><p><supplementary-material id="fig4sdata2"><label>Figure 4—source data 2.</label><caption><title><italic>C.</italic> <italic>difficile</italic> 20291 versus M7404, list of muropeptides, abundance, RT.</title></caption><media mime-subtype="xlsx" mimetype="application" xlink:href="elife-70597-fig4-data2-v1.xlsx"/></supplementary-material></p></caption><graphic mime-subtype="tiff" mimetype="image" xlink:href="elife-70597-fig4-v1.tif"/></fig><fig id="fig4s1" position="float" specific-use="child-fig"><label>Figure 4—figure supplement 1.</label><caption><title><italic>C.</italic> <italic>difficile</italic> LC-MS chromatograms.</title></caption><graphic mime-subtype="tiff" mimetype="image" xlink:href="elife-70597-fig4-figsupp1-v1.tif"/></fig></fig-group></sec><sec id="s2-4"><title>Benchmarking the automated PG analysis pipeline using available datasets</title><p>The analysis of <italic>E. coli</italic> PG established a proof of concept, showing that our matching script is suitable for the automated analysis of PG MS data. We next sought to evaluate the robustness of our peptidoglycomics pipeline using available datasets described in the literature. The most suitable publication was a recent study by Anderson et al. describing a PG analysis of <italic>P. aeruginosa</italic> planktonic cells (<xref ref-type="bibr" rid="bib2">Anderson et al., 2020b</xref>). Unlike most (if not all) studies published to date, this work provided datasets from biological and technical triplicates. Unlike our <italic>E. coli</italic> samples, analysed on a Q Exactive Focus Orbitrap (Thermo), <italic>P. aeruginosa</italic> samples were analysed on an Agilent Q-TOF mass spectrometer. Spectra were deconvoluted using the Byos Feature Finder module, and observed masses were matched using a mass tolerance of 25 ppm as described by Anderson et al. To limit the occurrence of mass coincidences and misidentification at this slightly lower mass accuracy, we carried out a search with PG modifications (anhydroMurNAc residues, deacetylation, lack of peptide stem resulting from amidase activity and amidation) but only using the most frequent combinations of modifications (double anhydroMurNAc and anhydroMurNAc and deacetyl).</p><p>PGFinder identified 63 muropeptides out of the 71 reported by Anderson et al., matching our search criteria (<xref ref-type="table" rid="table2">Table 2</xref>). The eight muropeptides that were not identified were absent from the list of deconvoluted masses, indicating that the problem was not associated with the script, highlighting that the deconvolution step is a source of variability. Interestingly, the observed masses calculated using Byos Feature Finder were closer to the theoretical value (6.5 ppm versus 10.7 ppm on average), reflecting another source of variability associated with data deconvolution. Four muropeptides previously identified containing the AEJAG pentapeptide stem were matched with distinct structures due to a mass coincidence between K and AG. A careful analysis of MS/MS spectra suggested that these muropeptides contained a K residue rather than the AG dipeptide (the y2 ion being 143 ppm away from the expected mass), showing the added value of an unbiased search. This conclusion is supported by the retention times of the corresponding muropeptides since the tetrapeptide AEJK elutes before EAJA whilst the pentapeptide AEJAG elutes later in the chromatography (<xref ref-type="bibr" rid="bib4">Bern et al., 2017</xref>). It is worth noting that our search also identified a large number of muropeptides that were not reported previously (<xref ref-type="supplementary-material" rid="table2sdata1">Table 2—source data 1</xref>). These results collectively show that our PG analysis pipeline can identify an unprecedentedly large number of muropeptide structures based on MS1 data, including all those previously reported (<xref ref-type="bibr" rid="bib1">Anderson et al., 2020a</xref>; <xref ref-type="bibr" rid="bib2">Anderson et al., 2020b</xref>).</p><table-wrap id="table2" position="float"><label>Table 2.</label><caption><title>Automated identification of <italic>P. aeruginosa</italic> peptidoglycan fragments.</title><p><supplementary-material id="table2sdata1"><label>Table 2—source data 1.</label><caption><title><italic>Pseudomonas aeruginosa</italic> matched muropeptides not reported previously.</title></caption><media mime-subtype="xlsx" mimetype="application" xlink:href="elife-70597-table2-data1-v1.xlsx"/></supplementary-material><supplementary-material id="table2sdata2"><label>Table 2—source data 2.</label><caption><title>Raw output of automated search using MaxQuant and PGFinder.</title></caption><media mime-subtype="xlsx" mimetype="application" xlink:href="elife-70597-table2-data2-v1.xlsx"/></supplementary-material></p></caption><table frame="hsides" rules="groups"><thead><tr><th align="left" rowspan="2" valign="bottom">Inferred structure</th><th align="left" colspan="2" valign="bottom">Mass</th><th align="left" colspan="2" valign="bottom">∆ppm</th><th align="left" rowspan="2" valign="bottom">MaxQuant</th></tr><tr><th align="left" valign="bottom">Theoretical</th><th align="left" valign="bottom">Observed</th><th align="left" valign="bottom">This work</th><th align="left" valign="bottom">Anderson et al.</th></tr></thead><tbody><tr><td align="left" valign="bottom">GM (anhydro)</td><td align="char" char="." valign="bottom">478.1799</td><td align="char" char="." valign="bottom">478.1780</td><td align="char" char="." valign="bottom">4.0</td><td align="char" char="." valign="bottom">–2.7</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM</td><td align="char" char="." valign="bottom">498.2061</td><td align="char" char="." valign="bottom">498.2042</td><td align="char" char="." valign="bottom">3.9</td><td align="char" char="." valign="bottom">–4.2</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM (x2) (deacetyl)</td><td align="char" char="." valign="bottom">934.3755</td><td align="char" char="." valign="bottom">934.3706</td><td align="char" char="." valign="bottom">5.3</td><td align="char" char="." valign="bottom">–8.6</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM (x2) (anhydro)</td><td align="char" char="." valign="bottom">956.3598</td><td align="char" char="." valign="bottom">956.3551</td><td align="char" char="." valign="bottom">5.0</td><td align="char" char="." valign="bottom">6.0</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM (x2)</td><td align="char" char="." valign="bottom">976.3860</td><td align="char" char="." valign="bottom">976.3794</td><td align="char" char="." valign="bottom">6.7</td><td align="char" char="." valign="bottom">–6.1</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM (x3) (deacetyl)</td><td align="char" char="." valign="bottom">1412.5554</td><td align="char" char="." valign="bottom">1412.5490</td><td align="char" char="." valign="bottom">4.5</td><td align="char" char="." valign="bottom">–6.2</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM (x3) (anhydro)</td><td align="char" char="." valign="bottom">1434.5397</td><td align="char" char="." valign="bottom">1434.5348</td><td align="char" char="." valign="bottom">3.4</td><td align="char" char="." valign="bottom">–7.5</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM (x3)</td><td align="char" char="." valign="bottom">1454.5659</td><td align="char" char="." valign="bottom">1454.5592</td><td align="char" char="." valign="bottom">4.6</td><td align="char" char="." valign="bottom">–5.3</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM (x4)</td><td align="char" char="." valign="bottom">1932.7458</td><td align="char" char="." valign="bottom">1932.7352</td><td align="char" char="." valign="bottom">5.5</td><td align="char" char="." valign="bottom">–5.1</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AE (anhydro)</td><td align="char" char="." valign="bottom">678.2596</td><td align="char" char="." valign="bottom">678.2567</td><td align="char" char="." valign="bottom">4.3</td><td align="char" char="." valign="bottom">–9.1</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AE</td><td align="char" char="." valign="bottom">698.2858</td><td align="char" char="." valign="bottom">698.2830</td><td align="char" char="." valign="bottom">3.9</td><td align="char" char="." valign="bottom">–12.9</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJ (anhydro)</td><td align="char" char="." valign="bottom">850.3444</td><td align="char" char="." valign="bottom">850.3401</td><td align="char" char="." valign="bottom">5.1</td><td align="char" char="." valign="bottom">–10.6</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJ</td><td align="char" char="." valign="bottom">870.3706</td><td align="char" char="." valign="bottom">870.3676</td><td align="char" char="." valign="bottom">3.5</td><td align="char" char="." valign="bottom">–5.9</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA (anhydro)</td><td align="char" char="." valign="bottom">921.3815</td><td align="char" char="." valign="bottom">921.3765</td><td align="char" char="." valign="bottom">5.4</td><td align="char" char="." valign="bottom">–9.9</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJG</td><td align="char" char="." valign="bottom">927.3920</td><td align="char" char="." valign="bottom">927.3868</td><td align="char" char="." valign="bottom">5.6</td><td align="char" char="." valign="bottom">–8.9</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA</td><td align="char" char="." valign="bottom">941.4077</td><td align="char" char="." valign="bottom">941.4045</td><td align="char" char="." valign="bottom">3.4</td><td align="char" char="." valign="bottom">–5.0</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJC</td><td align="char" char="." valign="bottom">973.3843</td><td align="char" char="." valign="bottom">973.3763</td><td align="char" char="." valign="bottom">8.2</td><td align="char" char="." valign="bottom">–2072.2</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJL</td><td align="char" char="." valign="bottom">983.4593</td><td align="char" char="." valign="bottom">983.4498</td><td align="char" char="." valign="bottom">9.6</td><td align="char" char="." valign="bottom">–15.5</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJK</td><td align="char" char="." valign="bottom">998.4703</td><td align="char" char="." valign="bottom">998.4624</td><td align="char" char="." valign="bottom">8.0</td><td align="char" char="." valign="bottom">–10.6</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJM</td><td align="char" char="." valign="bottom">1001.4153</td><td align="char" char="." valign="bottom">1001.4060</td><td align="char" char="." valign="bottom">9.2</td><td align="char" char="." valign="bottom">–13.5</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJAA</td><td align="char" char="." valign="bottom">1012.4448</td><td align="char" char="." valign="bottom">1012.4413</td><td align="char" char="." valign="bottom">3.4</td><td align="char" char="." valign="bottom">–7.8</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJY (anhydro)</td><td align="char" char="." valign="bottom">1013.4091</td><td align="char" char="." valign="bottom">1013.4242</td><td align="char" char="." valign="bottom">–14.9</td><td align="char" char="." valign="bottom">17.8</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJF</td><td align="char" char="." valign="bottom">1017.4433</td><td align="char" char="." valign="bottom">1017.4347</td><td align="char" char="." valign="bottom">8.4</td><td align="char" char="." valign="bottom">–15.0</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJY</td><td align="char" char="." valign="bottom">1033.4353</td><td align="char" char="." valign="bottom">1033.4278</td><td align="char" char="." valign="bottom">7.2</td><td align="char" char="." valign="bottom">–5.3</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJAV</td><td align="char" char="." valign="bottom">1040.4808</td><td align="char" char="." valign="bottom">1040.4716</td><td align="char" char="." valign="bottom">8.8</td><td align="char" char="." valign="bottom">–14.7</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJIA</td><td align="char" char="." valign="bottom">1054.4964</td><td align="char" char="." valign="bottom">1054.4874</td><td align="char" char="." valign="bottom">8.5</td><td align="char" char="." valign="bottom">–11.3</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJW</td><td align="char" char="." valign="bottom">1056.4394</td><td align="char" char="." valign="bottom">1056.4455</td><td align="char" char="." valign="bottom">–5.8</td><td align="char" char="." valign="bottom">4.0</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJAM</td><td align="char" char="." valign="bottom">1072.4524</td><td align="char" char="." valign="bottom">1072.4460</td><td align="char" char="." valign="bottom">5.9</td><td align="char" char="." valign="bottom">–4.3</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJKR</td><td align="char" char="." valign="bottom">1154.5667</td><td align="char" char="." valign="bottom">1154.5631</td><td align="char" char="." valign="bottom">3.1</td><td align="char" char="." valign="bottom">–8.1</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-GM-AE</td><td align="char" char="." valign="bottom">1176.4836</td><td align="char" char="." valign="bottom">1176.4590</td><td align="char" char="." valign="bottom">20.9</td><td align="char" char="." valign="bottom">–24.7</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-GM-AEJ</td><td align="char" char="." valign="bottom">1348.5684</td><td align="char" char="." valign="bottom">1348.5457</td><td align="char" char="." valign="bottom">16.9</td><td align="char" char="." valign="bottom">–24.9</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-GM-AEJA</td><td align="char" char="." valign="bottom">1419.6055</td><td align="char" char="." valign="bottom">1419.5824</td><td align="char" char="." valign="bottom">16.2</td><td align="char" char="." valign="bottom">–23.5</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA-GM-AEJ (amidase product)</td><td align="char" char="." valign="bottom">1313.5721</td><td align="char" char="." valign="bottom">1313.5674</td><td align="char" char="." valign="bottom">3.5</td><td align="char" char="." valign="bottom">–11.0</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA-GM-AEJA (amidase product)</td><td align="char" char="." valign="bottom">1384.6092</td><td align="char" char="." valign="bottom">1384.6037</td><td align="char" char="." valign="bottom">4.0</td><td align="char" char="." valign="bottom">–7.4</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJ-GM-AEJ (anhydro)</td><td align="char" char="." valign="bottom">1702.7042</td><td align="char" char="." valign="bottom">1702.6976</td><td align="char" char="." valign="bottom">3.9</td><td align="char" char="." valign="bottom">38.3</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJ-GM-AEJ</td><td align="char" char="." valign="bottom">1722.7304</td><td align="char" char="." valign="bottom">1722.7234</td><td align="char" char="." valign="bottom">4.1</td><td align="char" char="." valign="bottom">–8.6</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA-GM-AEJ (double anhydro)</td><td align="char" char="." valign="bottom">1753.7151</td><td align="char" char="." valign="bottom">1753.7043</td><td align="char" char="." valign="bottom">6.2</td><td align="char" char="." valign="bottom">–7.2</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA-GM-AEJ (anhydro)</td><td align="char" char="." valign="bottom">1773.7413</td><td align="char" char="." valign="bottom">1773.7339</td><td align="char" char="." valign="bottom">4.2</td><td align="char" char="." valign="bottom">–11.1</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA-GM-AEJ</td><td align="char" char="." valign="bottom">1793.7675</td><td align="char" char="." valign="bottom">1793.7596</td><td align="char" char="." valign="bottom">4.4</td><td align="char" char="." valign="bottom">–8.8</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA-GM-AEJA (dacetyl)</td><td align="char" char="." valign="bottom">1822.7941</td><td align="char" char="." valign="bottom">1822.7808</td><td align="char" char="." valign="bottom">7.3</td><td align="char" char="." valign="bottom">–7.4</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA-GM-AEJA (double anhydro)</td><td align="char" char="." valign="bottom">1824.7601</td><td align="char" char="." valign="bottom">1824.7447</td><td align="char" char="." valign="bottom">8.4</td><td align="char" char="." valign="bottom">–15.6</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA-GM-AEJA (anhydro)</td><td align="char" char="." valign="bottom">1844.7784</td><td align="char" char="." valign="bottom">1844.7704</td><td align="char" char="." valign="bottom">4.3</td><td align="char" char="." valign="bottom">–8.3</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA-GM-AEJG</td><td align="char" char="." valign="bottom">1850.7889</td><td align="char" char="." valign="bottom">1850.8158</td><td align="char" char="." valign="bottom">–14.6</td><td align="char" char="." valign="bottom">9.7</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA-GM-AEJA</td><td align="char" char="." valign="bottom">1864.8046</td><td align="char" char="." valign="bottom">1864.7962</td><td align="char" char="." valign="bottom">4.5</td><td align="char" char="." valign="bottom">–6.6</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA-GM-AEJK (anhydro)</td><td align="char" char="." valign="bottom">1901.8410</td><td align="char" char="." valign="bottom">1901.8297</td><td align="char" char="." valign="bottom">5.9</td><td align="char" char="." valign="bottom">–14.5</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA-GM-AEJL</td><td align="char" char="." valign="bottom">1906.8562</td><td align="char" char="." valign="bottom">1906.8452</td><td align="char" char="." valign="bottom">5.8</td><td align="char" char="." valign="bottom">–11.3</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA-GM-AEJK</td><td align="char" char="." valign="bottom">1921.8672</td><td align="char" char="." valign="bottom">1921.8586</td><td align="char" char="." valign="bottom">4.5</td><td align="char" char="." valign="bottom">–12.0</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA-GM-AEJF</td><td align="char" char="." valign="bottom">1940.8402</td><td align="char" char="." valign="bottom">1940.8263</td><td align="char" char="." valign="bottom">7.2</td><td align="char" char="." valign="bottom">–8.8</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA-GM-AEJY</td><td align="char" char="." valign="bottom">1956.8322</td><td align="char" char="." valign="bottom">1956.8210</td><td align="char" char="." valign="bottom">5.7</td><td align="char" char="." valign="bottom">–7.6</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA-GM-AEJAL</td><td align="char" char="." valign="bottom">1977.8933</td><td align="char" char="." valign="bottom">1977.8813</td><td align="char" char="." valign="bottom">6.0</td><td align="char" char="." valign="bottom">–10.7</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA-GM-AEJKR</td><td align="char" char="." valign="bottom">2077.9636</td><td align="char" char="." valign="bottom">2077.9589</td><td align="char" char="." valign="bottom">2.2</td><td align="char" char="." valign="bottom">–13.0</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-GM-AEJ-GM-AEJ</td><td align="char" char="." valign="bottom">2200.9282</td><td align="char" char="." valign="bottom">2200.9000</td><td align="char" char="." valign="bottom">12.8</td><td align="char" char="." valign="bottom">–17.7</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-GM-AEJA-GM-AEJ</td><td align="char" char="." valign="bottom">2271.9653</td><td align="char" char="." valign="bottom">2271.9368</td><td align="char" char="." valign="bottom">12.6</td><td align="char" char="." valign="bottom">–18.4</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-GM-AEJA-GM-AEJA</td><td align="char" char="." valign="bottom">2343.0024</td><td align="char" char="." valign="bottom">2342.9734</td><td align="char" char="." valign="bottom">12.4</td><td align="char" char="." valign="bottom">411.4</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA-GM-AEJA-GM-AEJ (double anhydro)</td><td align="char" char="." valign="bottom">2677.1120</td><td align="char" char="." valign="bottom">2677.1000</td><td align="char" char="." valign="bottom">4.5</td><td align="char" char="." valign="bottom">–10.7</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA-GM-AEJA-GM-AEJ (anhydro)</td><td align="char" char="." valign="bottom">2697.1382</td><td align="char" char="." valign="bottom">2697.1259</td><td align="char" char="." valign="bottom">4.6</td><td align="char" char="." valign="bottom">–8.6</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA-GM-AEJA-GM-AEJ</td><td align="char" char="." valign="bottom">2717.1644</td><td align="char" char="." valign="bottom">2717.1532</td><td align="char" char="." valign="bottom">4.1</td><td align="char" char="." valign="bottom">–10.7</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA-GM-AEJA-GM-AEJA (double anhydro)</td><td align="char" char="." valign="bottom">2748.1491</td><td align="char" char="." valign="bottom">2748.1363</td><td align="char" char="." valign="bottom">4.7</td><td align="char" char="." valign="bottom">–11.0</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA-GM-AEJA-GM-AEJA (anhydro)</td><td align="char" char="." valign="bottom">2768.1753</td><td align="char" char="." valign="bottom">2768.1674</td><td align="char" char="." valign="bottom">2.9</td><td align="char" char="." valign="bottom">–11.2</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA-GM-AEJA-GM-AEJA</td><td align="char" char="." valign="bottom">2788.2015</td><td align="char" char="." valign="bottom">2788.1919</td><td align="char" char="." valign="bottom">3.4</td><td align="char" char="." valign="bottom">–9.7</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-AEJA-GM-AEJA-GM-AEJK (anhydro)</td><td align="char" char="." valign="bottom">2825.2379</td><td align="char" char="." valign="bottom">2825.2205</td><td align="char" char="." valign="bottom">6.1</td><td align="char" char="." valign="bottom">–9.3</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-GM-AEJA-GM-AEJA-GM-AEJ</td><td align="char" char="." valign="bottom">3195.3622</td><td align="char" char="." valign="bottom">3195.3264</td><td align="char" char="." valign="bottom">11.2</td><td align="char" char="." valign="bottom">–14.0</td><td align="char" char="." valign="bottom">+</td></tr><tr><td align="left" valign="bottom">GM-GM-AEJA-GM-AEJA-GM-AEJA</td><td align="char" char="." valign="bottom">3266.3993</td><td align="char" char="." valign="bottom">3266.3630</td><td align="char" char="." valign="bottom">11.1</td><td align="char" char="." valign="bottom">–12.5</td><td align="char" char="." valign="bottom">+</td></tr></tbody></table><table-wrap-foot><fn><p>Alternative structures were matched:</p></fn><fn><p>GM-AEJ-GM-AEJK.</p></fn><fn><p>GM-AEJ-GM-AEJKA (anhydro).</p></fn><fn><p>GM-AEJ-GM-AEJKA.</p></fn><fn><p>GM-AEJ-GM-AEJA-GM-AEJKA (anhydro).</p></fn></table-wrap-foot></table-wrap></sec><sec id="s2-5"><title>A PG analysis workflow using freely available tools</title><p>The automated identification of <italic>P. aeruginosa</italic> muropeptides using PGFinder and the strategy reported previously (<xref ref-type="bibr" rid="bib1">Anderson et al., 2020a</xref>; <xref ref-type="bibr" rid="bib2">Anderson et al., 2020b</xref>) relied on commercially available deconvolution software (ProteinMetrics Byos or Agilent MassHunter, respectively). We sought to identify a free alternative software for mass deconvolution to make our PG analysis pipeline accessible to everyone. MaxQuant was selected as a tool of choice since it represents a widely used software package to analyse high-resolution mass-spectrometric data for shotgun proteomics (<xref ref-type="bibr" rid="bib6">Cox and Mann, 2008</xref>). MaxQuant has native support for Thermo MS data (RAW) and Sciex (WIFF) file formats and supports the open MS data format mzXML. Virtually any proprietary MS data file can be converted to the mzXML open format using freely available tools such as Proteowizard and TOPPAS, making this workflow universally applicable. As proof of concept, we converted <italic>P. aeruginosa</italic> data to an mzXML file (<xref ref-type="fig" rid="app2fig1">Appendix 2—figure 1</xref>), processed it using MaxQuant for mass deconvolution (<xref ref-type="fig" rid="app2fig2">Appendix 2—figure 2</xref>), and analysed it using PGFinder for PG structure and composition identification. We were able to identify all the expected muropeptides (<xref ref-type="table" rid="table2">Table 2</xref>). This result confirms that the automated analysis of PG datasets can be carried out using the MaxQuant freeware and our open-source script PGFinder (<xref ref-type="table" rid="table2">Table 2</xref> and <xref ref-type="supplementary-material" rid="table2sdata2">Table 2—source data 2</xref>).</p></sec></sec><sec id="s3" sec-type="discussion"><title>Discussion</title><p>This study describes a workflow for the unbiased and automated analysis of bacterial PG using freely available resources. We analysed high-resolution MS datasets corresponding to PG fragments and demonstrated that this approach is a powerful tool to identify muropeptides and carry out comparative analyses based on the MS1 data.</p><p>MS analysis of bacterial PG has been carried out since the late 1980s (<xref ref-type="bibr" rid="bib10">Garcia-Bustos et al., 1988</xref>; <xref ref-type="bibr" rid="bib18">Martin et al., 1987</xref>). Whilst rp-HPLC-MS is now routinely used to explore PG structure, data analysis remains a manual process. Therefore, this step has become a bottleneck that prevents high-throughput analyses and introduces a series of issues regarding reproducibility. A major issue deals with the definition of the search space used since the list of muropeptides searched is not provided. Our previous work showed that an unbiased approach using shotgun proteomics tools identifies muropeptides containing unusual amino acids in <italic>C. difficile</italic> (<xref ref-type="bibr" rid="bib4">Bern et al., 2017</xref>). Our unbiased search identified &gt;106 masses matching <italic>E. coli</italic> muropeptide structures, representing a number strikingly larger than previously reported (<xref ref-type="bibr" rid="bib15">Kühner et al., 2014</xref>). This level of complexity was anticipated based on the complement of enzymes involved in <italic>E. coli</italic> PG synthesis but had never been reported before. Our work therefore highlighted the limitations of search strategies reported so far, especially if we consider the fact that some bacterial PGs like <italic>E. coli</italic> have been extensively studied over the past 30 years.</p><p>Several issues and flaws associated with the manual analysis of PG MS data are addressed by our approach, which represents a robust, consistent, and open-access strategy. The PGFinder algorithm has been designed to create dynamic databases that are ultimately combined to perform the final matching process. Optimisation of the search space relies on a preliminary identification of masses matching theoretical monoisotopic masses of monomers, limiting misidentifications based on mass coincidence. Another advantage is that the only information provided by the user is a restricted monomer database rather than a comprehensive one. This avoids a time-consuming operation, prone to human error. PGFinder uses XIC for quantification of PG fragments, providing high resolution, sensitivity, and reproducibility, as indicated by the comparisons across biological replicates (<xref ref-type="fig" rid="app1fig2">Appendix 1—figure 2</xref> and <xref ref-type="fig" rid="fig4">Figure 4a</xref>). It enables the accurate quantification of molecules with overlapping retention times and those present in very low abundance, with a large dynamic range (typically six orders of magnitude). Although a Matlab-based software package (Chromanalysis) has been described to automate the detection and quantification of UV peaks through Gaussian fitting (<xref ref-type="bibr" rid="bib8">Desmarais et al., 2015</xref>), quantification using XIC is more straightforward. It is worth pointing out that the approach described here does not give absolute quantification of muropeptides. Instead, it allows a relative quantification of muropeptides. Nevertheless, this strategy remains suitable for comparative analyses and overcomes two major limitations associated with UV detection, namely detection threshold and co-elution of molecules.</p><p>Our <italic>E. coli</italic> PG analysis confirmed that PGFinder is a powerful tool that provides a much improved qualitative and quantitative PG analysis (<xref ref-type="bibr" rid="bib15">Kühner et al., 2014</xref>; <xref ref-type="bibr" rid="bib19">Morè et al., 2019</xref>). Combining an unbiased search with highly sensitive detection of individual structures is important for two reasons. Firstly, it opens the possibility to identify subtle modifications of PG structure, resulting from either a transient or a localised enzymatic activity such as that taking place at the septum. Secondly, it will permit the identification of previously undetected modifications that may provide new insights into our understanding of PG composition and dynamics. For example, we showed that <italic>E. coli</italic> PG contains a low abundance of deacetylated sugars. This observation is puzzling because no canonical PG deacetylase genes have been identified in this organism. Although the biological relevance of this property remains to be established, we cannot exclude the possibility that PG deacetylation in <italic>E. coli</italic> may contribute to PG homeostasis. Another striking outcome resulting from our automated search is the identification of a slightly higher amount of muropeptides containing anhydromuramic acid (4.55%) as compared to 2–3% (<xref ref-type="bibr" rid="bib12">Glauner et al., 1988</xref>; <xref ref-type="bibr" rid="bib16">Liu et al., 2020</xref>). To explain the discrepancy between our work and data from the literature, it is tempting to assume that most of the muropeptides containing anhydromuramic acid identified with PGFinder were simply not searched in previous studies. It is worth pointing out that none of the papers describing PG analysis published to date has reported the list of structures searched in the MS data analysed.</p><p>One of our objectives was to create an automated PG analysis tool accessible to the broadest audience possible, including people with no prior experience with programming or coding languages. Therefore, we shared PGFinder as a Jupyter Notebook allowing users to customise the search strategy depending on both the question asked and the instrument accuracy. PGFinder is particularly suitable for the characterisation of novel PGs with unknown composition or structural modifications and can be modified by users to add novel functionalities. However, a current limitation of this workflow is that it does not process MS/MS data. Therefore, the fragmentation spectra of individual monomers must be checked using dedicated tools to validate that the inferred structures are correct. We are currently working towards an integrated pipeline that includes MS/MS analysis to our PGFinder pipeline. The ability to disable some PG modifications means that the complexity of the search can be adjusted to focus on specific properties (e.g., the occurrence of acetylation/deacetylation, or amidation) or specific muropeptides resulting from lytic activities (e.g., unsubstituted MurNAc residues resulting from amidase activity). For PG that have already been well characterised (<italic>E. coli</italic>, <italic>P. aeruginosa,</italic> or <italic>C. difficile</italic>), the search parameters are already established, allowing a very straightforward analysis to be performed. Therefore, access to a custom, semi-quantitative sensitive analysis is ideal for comparing PG dynamics or differences in PG structure between a reference strain and isogenic mutants. Both reduced disaccharide peptides or lactyl-peptides (generated by beta-elimination) are identified using PGFinder.</p><p>We anticipate that an open access to PGFinder, in conjunction with freely available deconvolution tools, will allow researchers to carry out comparative MS1 analyses. The pipeline defined in this work enables reproducible and consistent data analysis. This represents the first step towards a standardised approach to PG analysis, opening the possibility to reanalyse datasets in repositories. The modular structure of the open-source PGFinder code can be easily integrated into any specific workflow for the automated processing of PG MS data.</p></sec><sec id="s4" sec-type="materials|methods"><title>Materials and methods</title><table-wrap id="keyresource" position="anchor"><label>Key resources table</label><table frame="hsides" rules="groups"><thead><tr><th align="left" valign="bottom">Reagent type (species) or resource</th><th align="left" valign="bottom">Designation</th><th align="left" valign="bottom">Source or reference</th><th align="left" valign="bottom">Identifiers</th><th align="left" valign="bottom">Additional information</th></tr></thead><tbody><tr><td align="left" valign="bottom">Strain, strain background(<italic>Escherichia coli</italic>)</td><td align="left" valign="bottom">BW25113</td><td align="left" valign="bottom"><ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1073/pnas.120163297">https://doi.org/10.1073/pnas.120163297</ext-link></td><td align="left" valign="bottom">RRID:<ext-link ext-link-type="uri" xlink:href="https://identifiers.org/RRID/RRID:Addgene_72340">Addgene_72340</ext-link></td><td align="left" valign="bottom">Model strain for PG analysis</td></tr><tr><td align="left" valign="bottom">Strain, strain background(<italic>Clostridioides difficile</italic>)</td><td align="left" valign="bottom">R20291</td><td align="left" valign="bottom"><ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1128/JB.0073107">https://doi.org/10.1128/JB.0073107</ext-link></td><td align="left" valign="bottom"/><td align="left" valign="bottom">Model strain for PG analysis</td></tr><tr><td align="left" valign="bottom">Strain, strain background(<italic>Clostridioides difficile</italic>)</td><td align="left" valign="bottom">M7404</td><td align="left" valign="bottom"><ext-link ext-link-type="uri" xlink:href="https://doi.org/10.1371/journal.ppat.1002317">https://doi.org/10.1371/journal.ppat.1002317</ext-link></td><td align="left" valign="bottom"/><td align="left" valign="bottom">Model strain for PG analysis</td></tr><tr><td align="left" valign="bottom">Software, algorithm</td><td align="left" valign="bottom">PGFinderv.0.02</td><td align="left" valign="bottom">This work</td><td align="left" valign="bottom"/><td align="left" valign="bottom">Used for MS1 analysis of PG structure</td></tr><tr><td align="left" valign="bottom">Software, algorithm</td><td align="left" valign="bottom">Byosv.3.9–32</td><td align="left" valign="bottom">Protein Metrics Inc</td><td align="left" valign="bottom"/><td align="left" valign="bottom">Used for MS data deconvolution and MS/MS analysis</td></tr><tr><td align="left" valign="bottom">Software, algorithm</td><td align="left" valign="bottom">MaxQuant v2.0.1.0</td><td align="left" valign="bottom"><xref ref-type="bibr" rid="bib6">Cox and Mann, 2008</xref></td><td align="left" valign="bottom">RRID:<ext-link ext-link-type="uri" xlink:href="https://identifiers.org/RRID/RRID:SCR_014485">SCR_014485</ext-link></td><td align="left" valign="bottom">Used for MS data deconvolution</td></tr><tr><td align="left" valign="bottom">Software, algorithm</td><td align="left" valign="bottom">Perseusv.1.6.10.53</td><td align="left" valign="bottom"><xref ref-type="bibr" rid="bib26">Tyanova et al., 2016</xref></td><td align="left" valign="bottom">RRID:<ext-link ext-link-type="uri" xlink:href="https://identifiers.org/RRID/RRID:SCR_015753">SCR_015753</ext-link></td><td align="left" valign="bottom">Used statistical analysis of muropeptide abundance</td></tr></tbody></table></table-wrap><sec id="s4-1"><title>Bacterial strains and culture conditions</title><p><italic>E. coli</italic> BW25113 was grown at 37°C in LB under agitation (250 rpm). <italic>C. difficile</italic> strains were cultured in heart infusion supplemented with yeast extract, L-cysteine, and glucose in an atmosphere of 10% H<sub>2</sub>, 10% CO<sub>2</sub>, and 80% N<sub>2</sub> at 37°C in a Coy chamber or Don Whitley A300 anaerobic workstation.</p></sec><sec id="s4-2"><title>PG purification</title><p>PG was purified from exponential (<italic>E. coli</italic>) or late exponential (<italic>C. difficile</italic>) phase as described previously (<xref ref-type="bibr" rid="bib9">Eckert et al., 2006</xref>; <xref ref-type="bibr" rid="bib11">Glauner, 1988</xref>), freeze-dried, and resuspended in distilled water at a concentration of 5 mg/ml.</p></sec><sec id="s4-3"><title>Preparation of soluble muropeptides</title><p>PG (1 mg) was digested overnight with 25 µg of mutanolysin at 37°C in 150 µl of 20 mM sodium phosphate buffer (pH 5.5). Soluble disaccharide peptides were recovered in the supernatant following centrifugation (20,000 × <italic>g</italic> for 20 min at 25°C). To reduce muropeptides, equal volumes (200 µl) of the solution of disaccharide peptides and of borate buffer (250 mM, pH 9.0) were mixed. 2 ml of sodium borohydride was added, and the solution was incubated for 20 min at room temperature. The pH of the solution was adjusted to 4.0 with 20% orthophosphoric acid. Beta-elimination was carried out by mixing 200 µl of muropeptides with 64 µl of 32% (w/v) ammonia. After 5 hr at 37 C, the solution was neutralised with 60 µl of acetic acid glacial, freeze-dried, and resuspended in water (<xref ref-type="bibr" rid="bib3">Arbeloa et al., 2004</xref>; <xref ref-type="bibr" rid="bib9">Eckert et al., 2006</xref>).</p><p>The reduced muropeptides were desalted by reverse-phase high-performance liquid chromatography (rp-HPLC) on a C18 Hypersil Gold aQ column (3 µm, 2.1 × 200 mm; Thermo Fisher) at a flow rate of 0.4 ml/min. After 1 min in water-0.1% formic acid (v/v) (buffer A), muropeptides were eluted with a 6 min linear gradient to 95% acetonitrile-0.1% formic acid (v/v). Muropeptides were freeze-dried and resuspended in 100 µl. An aliquot of the desalted samples was analysed by rp-HPLC on the same column to measure the UV absorbance of the most abundant monomer (no isocratic step, muropeptides were eluted with a 30 min linear gradient to 15% acetonitrile-0.1% formic acid [v/v]). Samples were diluted to contain 150 mAU/µl of the major monomer and 10 µl were injected. Based on the dry weight of the PG sample, we estimated that this corresponded to approximately 50 µg of material.</p></sec><sec id="s4-4"><title>UHPLC-MS/MS</title><p>An Ultimate 3000 Ultra High-Performance Chromatography (UHPLC; Dionex/Thermo Fisher Scientific) system coupled with a high-resolution Q Exactive Focus mass spectrometer (Thermo Fisher Scientific) was used for LC/HRMS analysis. Muropeptides were separated using a C18 analytical column (Hypersil Gold aQ, 1.9 µm particles, 150 × 2.1 mm; Thermo Fisher Scientific), column temperature at 50°C. Muropeptides elution was performed by applying a mixture of solvent A (water, 0.1% [v/v] formic acid) and solvent B (acetonitrile, 0.1% [v/v] formic acid). After 10 µl sample injection, MS/MS data were acquired during a 40 min step gradient: 0–12.5% B for 25 min; 12.5–20% B for 5 min; held at 20% B for 5 min, and the column was re-equilibrated for 10 min under the initial conditions.</p><p>The Q Exactive Focus was operated under electrospray ionization (H-ESI II)-positive mode. Full scan (<italic>m/z</italic> 150–2250) used resolution 70,000 (FWHM) at <italic>m/z</italic> 200, with an automatic gain control (AGC) target of 1 × 10<sup>6</sup> ions and an automated maximum ion injection time (IT).</p><p>Data-dependent MS/MS were acquired on a ‘Top 3’ data-dependent mode using the following parameters: resolution 17,500; AGC 1 × 10<sup>5</sup> ions, maximum IT 50 ms, NCE 25%, and a dynamic exclusion time 5 s.</p></sec><sec id="s4-5"><title>MS data deconvolution</title><p>Byos search parameters to get .ftrs file Protein Metrics Byos v.3.9–32 was used to identify and compute the XICs. The parameters used for mass deconvolution using MaxQuant v.2.0.1.0 are described in <xref ref-type="fig" rid="app2fig2">Appendix 2—figure 2</xref>.</p><sec id="s4-5-1"><title>Data analysis</title><p>The crosslinking index and glycan chain length were calculated as described previously (<xref ref-type="bibr" rid="bib11">Glauner, 1988</xref>). Label-free relative quantitation of muropeptides from triplicate <italic>C. difficile</italic> clinical isolates (R20291 and M7404) was performed using Byos 3.11, and statistical analysis of the quantitative data was performed using Perseus v. 1.6.10.53 (<xref ref-type="bibr" rid="bib26">Tyanova et al., 2016</xref>). Briefly, muropeptide intensities were log<sub>2</sub> transformed and normalised by subtraction of the median value. A two-sample Student’s <italic>t</italic>-test was performed with a permutation-based FDR of 0.05 to determine statistically significant quantitative differences between the strains. Comparisons between R20291 and M7404 muropeptide distribution (mono-, di-, tri-, tetra-pentamer) was evaluated for statistical significance using GraphPad Prism (unpaired <italic>t</italic>-test).</p></sec></sec><sec id="s4-6"><title>Runtime environment</title><p>Code is available at <ext-link ext-link-type="uri" xlink:href="https://github.com/Mesnage-Org/PGFinder">https://github.com/Mesnage-Org/PGFinder</ext-link> (<xref ref-type="bibr" rid="bib21">Patel, 2021</xref>). <ext-link ext-link-type="uri" xlink:href="https://github.com/Mesnage-Org/pgfinder/releases/tag/v0.02">https://github.com/Mesnage-Org/PGFinder/releases/tag/v0.02</ext-link> will take you to the archived release used in this paper. We used Python 3 to write the MS1 package and demonstrate its functionality using demo scripts. PGFinder can be run through an interactive Jupyter Notebook hosted on mybinder for ease of use by those less familiar with Python code. A conda environment is provided to ensure reproducible execution. Regression testing has been implemented to ensure changes to code do not cause changes to important results. The GitHub contains an interactive version to run user’s analysis and an end-to-end demo using samples data provided with the script (<ext-link ext-link-type="uri" xlink:href="https://mybinder.org/v2/gh/Mesnage-Org/PGFinder/master?urlpath=tree/pgfinder_interactive.ipynb">Interactive PGFinder</ext-link>). The sample data is a MaxQuant deconvolution output from the <italic>E. coli</italic> MS data analysed in the paper. The current version of the script can handle both .txt (MaxQuant) or .ftrs (Byos) deconvoluted data and offers the possibility for the user to include several modifications in the search. The time window for the ‘clean up step’ (in-source decay and salt adducts) as well as ppm tolerance for matching can also be defined by the user; the default values corresponding to these parameters used in this work are 0.5 min and 10 ppm.</p></sec><sec id="s4-7"><title>Data availability</title><p>All <italic>E. coli</italic> and <italic>C. difficile</italic> MS datasets generated in this study are available through the GlycoPOST repository (GPST000168; <xref ref-type="bibr" rid="bib28">Watanabe et al., 2021</xref>). <italic>P. aeruginosa</italic> MS datasets are accessible via Figshare (<xref ref-type="bibr" rid="bib2">Anderson et al., 2020b</xref>).</p></sec></sec></body><back><sec id="s5" sec-type="additional-information"><title>Additional information</title><fn-group content-type="competing-interest"><title>Competing interests</title><fn fn-type="COI-statement" id="conf1"><p>none</p></fn><fn fn-type="COI-statement" id="conf2"><p>None</p></fn><fn fn-type="COI-statement" id="conf3"><p>Andrew Nichols is affiliated with Protein Metrics Inc. The author has no other competing interests to declare</p></fn><fn fn-type="COI-statement" id="conf4"><p>Marshall Bern is affiliated with Protein Metrics Inc. The author has no other competing interests to declare</p></fn></fn-group><fn-group content-type="author-contribution"><title>Author contributions</title><fn fn-type="con" id="con1"><p>Conceptualization, Data curation, Formal analysis, Investigation, Methodology, Resources, Software, Validation, Visualization, Writing – original draft, Writing – review and editing</p></fn><fn fn-type="con" id="con2"><p>Conceptualization, Data curation, Methodology, Software, Validation, Visualization, Writing – review and editing</p></fn><fn fn-type="con" id="con3"><p>Investigation, Methodology, Writing – review and editing</p></fn><fn fn-type="con" id="con4"><p>Data curation, Formal analysis, Investigation, Methodology, Resources, Writing – review and editing</p></fn><fn fn-type="con" id="con5"><p>Methodology, Resources, Funding acquisition</p></fn><fn fn-type="con" id="con6"><p>Project administration, Investigation, Funding acquisition, Visualization, Writing – review and editing</p></fn><fn fn-type="con" id="con7"><p>Supervision, Funding acquisition, Visualization, Writing – review and editing</p></fn><fn fn-type="con" id="con8"><p>Project administration, Methodology, Resources, Writing – review and editing</p></fn><fn fn-type="con" id="con9"><p>Conceptualization, Project administration, Methodology, Supervision, Software, Funding acquisition, Validation, Writing – review and editing</p></fn><fn fn-type="con" id="con10"><p>Formal analysis, Project administration, Methodology, Supervision, Resources, Funding acquisition, Validation, Visualization, Writing – review and editing</p></fn><fn fn-type="con" id="con11"><p>Conceptualization, Data curation, Formal analysis, Project administration, Investigation, Methodology, Supervision, Resources, Funding acquisition, Validation, Visualization, Writing – original draft, Writing – review and editing</p></fn></fn-group></sec><sec id="s6" sec-type="supplementary-material"><title>Additional files</title><supplementary-material id="transrepform"><label>Transparent reporting form</label><media mime-subtype="docx" mimetype="application" xlink:href="elife-70597-transrepform1-v1.docx"/></supplementary-material><supplementary-material id="supp1"><label>Supplementary file 1.</label><caption><title>Step by step strategy for PG analysis.</title></caption><media mime-subtype="pdf" mimetype="application" xlink:href="elife-70597-supp1-v1.pdf"/></supplementary-material></sec><sec id="s7" sec-type="data-availability"><title>Data availability</title><p>All raw mass spectrometry data files have are available through the Glycopost repository (ref GPST000168).</p><p>The following previously published datasets were used:</p><p><element-citation id="dataset1" publication-type="data" specific-use="references"><person-group person-group-type="author"><name><surname>Anderson</surname><given-names>EM</given-names></name><name><surname>Sychantha</surname><given-names>D</given-names></name><name><surname>Brewer</surname><given-names>D</given-names></name><name><surname>Clarke</surname><given-names>AJ</given-names></name><name><surname>Geddes-McAlister</surname><given-names>J</given-names></name><name><surname>Khursigara</surname><given-names>CM</given-names></name></person-group><year iso-8601-date="2020">2020</year><data-title>Peptidoglycomics: Examining compositional changes in peptidoglycan between biofilm- and planktonic-derived <italic>Pseudomonas aeruginosa</italic></data-title><source>figshare</source><pub-id pub-id-type="doi">10.6084/m9.figshare.10277909</pub-id></element-citation></p></sec><ack id="ack"><title>Acknowledgements</title><p>We thank Dominique Mengin-Lecreulx (University Paris XI) for <italic>E. coli</italic> strain BW25113. Yong Kil and Eric Carslon (Protein Metrics) are acknowledged for their constant support. Motoshi Suzuki (NIH/NIAID) is acknowledged for insightful discussions.</p></ack><ref-list><title>References</title><ref id="bib1"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Anderson</surname><given-names>EM</given-names></name><name><surname>Greenwood</surname><given-names>NA</given-names></name><name><surname>Brewer</surname><given-names>D</given-names></name><name><surname>Khursigara</surname><given-names>CM</given-names></name></person-group><year iso-8601-date="2020">2020a</year><article-title>Semi-quantitative analysis of peptidoglycan by liquid chromatography mass spectrometry and bioinformatics</article-title><source>Journal of Visualized Experiments</source><volume>164</volume><elocation-id>8</elocation-id><pub-id pub-id-type="doi">10.3791/61799</pub-id></element-citation></ref><ref id="bib2"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Anderson</surname><given-names>EM</given-names></name><name><surname>Sychantha</surname><given-names>D</given-names></name><name><surname>Brewer</surname><given-names>D</given-names></name><name><surname>Clarke</surname><given-names>AJ</given-names></name><name><surname>Geddes-McAlister</surname><given-names>J</given-names></name><name><surname>Khursigara</surname><given-names>CM</given-names></name></person-group><year iso-8601-date="2020">2020b</year><article-title>Peptidoglycomics reveals compositional changes in peptidoglycan between biofilm- and planktonic-derived <italic>Pseudomonas aeruginosa</italic></article-title><source>The Journal of Biological Chemistry</source><volume>295</volume><fpage>504</fpage><lpage>516</lpage><pub-id pub-id-type="doi">10.1074/jbc.RA119.010505</pub-id><pub-id pub-id-type="pmid">31771981</pub-id></element-citation></ref><ref id="bib3"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Arbeloa</surname><given-names>A</given-names></name><name><surname>Hugonnet</surname><given-names>JE</given-names></name><name><surname>Sentilhes</surname><given-names>AC</given-names></name><name><surname>Josseaume</surname><given-names>N</given-names></name><name><surname>Dubost</surname><given-names>L</given-names></name><name><surname>Monsempes</surname><given-names>C</given-names></name><name><surname>Blanot</surname><given-names>D</given-names></name><name><surname>Brouard</surname><given-names>JP</given-names></name><name><surname>Arthur</surname><given-names>M</given-names></name></person-group><year iso-8601-date="2004">2004</year><article-title>Synthesis of mosaic peptidoglycan cross-bridges by hybrid peptidoglycan assembly pathways in gram-positive bacteria</article-title><source>J Biol Chem</source><volume>279</volume><fpage>41546</fpage><lpage>41556</lpage><pub-id pub-id-type="doi">10.1074/jbc.M407149200</pub-id><pub-id pub-id-type="pmid">15280360</pub-id></element-citation></ref><ref id="bib4"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Bern</surname><given-names>M</given-names></name><name><surname>Beniston</surname><given-names>R</given-names></name><name><surname>Mesnage</surname><given-names>S</given-names></name></person-group><year iso-8601-date="2017">2017</year><article-title>Towards an automated analysis of bacterial peptidoglycan structure</article-title><source>Anal Bioanal Chem</source><volume>409</volume><fpage>551</fpage><lpage>560</lpage><pub-id pub-id-type="doi">10.1007/s00216-016-9857-5</pub-id><pub-id pub-id-type="pmid">27520322</pub-id></element-citation></ref><ref id="bib5"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Boneca</surname><given-names>IG</given-names></name><name><surname>Dussurget</surname><given-names>O</given-names></name><name><surname>Cabanes</surname><given-names>D</given-names></name><name><surname>Nahori</surname><given-names>M-A</given-names></name><name><surname>Sousa</surname><given-names>S</given-names></name><name><surname>Lecuit</surname><given-names>M</given-names></name><name><surname>Psylinakis</surname><given-names>E</given-names></name><name><surname>Bouriotis</surname><given-names>V</given-names></name><name><surname>Hugot</surname><given-names>J-P</given-names></name><name><surname>Giovannini</surname><given-names>M</given-names></name><name><surname>Coyle</surname><given-names>A</given-names></name><name><surname>Bertin</surname><given-names>J</given-names></name><name><surname>Namane</surname><given-names>A</given-names></name><name><surname>Rousselle</surname><given-names>J-C</given-names></name><name><surname>Cayet</surname><given-names>N</given-names></name><name><surname>Prévost</surname><given-names>M-C</given-names></name><name><surname>Balloy</surname><given-names>V</given-names></name><name><surname>Chignard</surname><given-names>M</given-names></name><name><surname>Philpott</surname><given-names>DJ</given-names></name><name><surname>Cossart</surname><given-names>P</given-names></name><name><surname>Girardin</surname><given-names>SE</given-names></name></person-group><year iso-8601-date="2007">2007</year><article-title>A critical role for peptidoglycan N-deacetylation in <italic>Listeria</italic> evasion from the host innate immune system</article-title><source>PNAS</source><volume>104</volume><fpage>997</fpage><lpage>1002</lpage><pub-id pub-id-type="doi">10.1073/pnas.0609672104</pub-id><pub-id pub-id-type="pmid">17215377</pub-id></element-citation></ref><ref id="bib6"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Cox</surname><given-names>J</given-names></name><name><surname>Mann</surname><given-names>M</given-names></name></person-group><year iso-8601-date="2008">2008</year><article-title>MaxQuant enables high peptide identification rates, individualized p.p.b.-range mass accuracies and proteome-wide protein quantification</article-title><source>Nature Biotechnology</source><volume>26</volume><fpage>1367</fpage><lpage>1372</lpage><pub-id pub-id-type="doi">10.1038/nbt.1511</pub-id><pub-id pub-id-type="pmid">19029910</pub-id></element-citation></ref><ref id="bib7"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Cummins</surname><given-names>CS</given-names></name><name><surname>Harris</surname><given-names>H</given-names></name></person-group><year iso-8601-date="1956">1956</year><article-title>The chemical composition of the cell wall in some gram-positive bacteria and its possible value as a taxonomic character</article-title><source>Journal of General Microbiology</source><volume>14</volume><fpage>583</fpage><lpage>600</lpage><pub-id pub-id-type="doi">10.1099/00221287-14-3-583</pub-id><pub-id pub-id-type="pmid">13346020</pub-id></element-citation></ref><ref id="bib8"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Desmarais</surname><given-names>SM</given-names></name><name><surname>Tropini</surname><given-names>C</given-names></name><name><surname>Miguel</surname><given-names>A</given-names></name><name><surname>Cava</surname><given-names>F</given-names></name><name><surname>Monds</surname><given-names>RD</given-names></name><name><surname>de Pedro</surname><given-names>MA</given-names></name><name><surname>Huang</surname><given-names>KC</given-names></name></person-group><year iso-8601-date="2015">2015</year><article-title>High-throughput, highly sensitive analyses of bacterial morphogenesis using ultra performance liquid chromatography</article-title><source>The Journal of Biological Chemistry</source><volume>290</volume><fpage>31090</fpage><lpage>31100</lpage><pub-id pub-id-type="doi">10.1074/jbc.M115.661660</pub-id><pub-id pub-id-type="pmid">26468288</pub-id></element-citation></ref><ref id="bib9"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Eckert</surname><given-names>C</given-names></name><name><surname>Lecerf</surname><given-names>M</given-names></name><name><surname>Dubost</surname><given-names>L</given-names></name><name><surname>Arthur</surname><given-names>M</given-names></name><name><surname>Mesnage</surname><given-names>S</given-names></name></person-group><year iso-8601-date="2006">2006</year><article-title>Functional analysis of AtlA, the major N-acetylglucosaminidase of <italic>Enterococcus faecalis</italic></article-title><source>Journal of Bacteriology</source><volume>188</volume><fpage>8513</fpage><lpage>8519</lpage><pub-id pub-id-type="doi">10.1128/JB.01145-06</pub-id><pub-id pub-id-type="pmid">17041059</pub-id></element-citation></ref><ref id="bib10"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Garcia-Bustos</surname><given-names>JF</given-names></name><name><surname>Chait</surname><given-names>BT</given-names></name><name><surname>Tomasz</surname><given-names>A</given-names></name></person-group><year iso-8601-date="1988">1988</year><article-title>Altered peptidoglycan structure in a pneumococcal transformant resistant to penicillin</article-title><source>Journal of Bacteriology</source><volume>170</volume><fpage>2143</fpage><lpage>2147</lpage><pub-id pub-id-type="doi">10.1128/jb.170.5.2143-2147.1988</pub-id><pub-id pub-id-type="pmid">3360741</pub-id></element-citation></ref><ref id="bib11"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Glauner</surname><given-names>B</given-names></name></person-group><year iso-8601-date="1988">1988</year><article-title>Separation and quantification of muropeptides with high-performance liquid chromatography</article-title><source>Analytical Biochemistry</source><volume>172</volume><fpage>451</fpage><lpage>464</lpage><pub-id pub-id-type="doi">10.1016/0003-2697(88)90468-x</pub-id><pub-id pub-id-type="pmid">3056100</pub-id></element-citation></ref><ref id="bib12"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Glauner</surname><given-names>B</given-names></name><name><surname>Höltje</surname><given-names>JV</given-names></name><name><surname>Schwarz</surname><given-names>U</given-names></name></person-group><year iso-8601-date="1988">1988</year><article-title>The composition of the murein of <italic>Escherichia coli</italic></article-title><source>The Journal of Biological Chemistry</source><volume>263</volume><fpage>10088</fpage><lpage>10095</lpage></element-citation></ref><ref id="bib13"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Harz</surname><given-names>H</given-names></name><name><surname>Burgdorf</surname><given-names>K</given-names></name><name><surname>Höltje</surname><given-names>JV</given-names></name></person-group><year iso-8601-date="1990">1990</year><article-title>Isolation and separation of the glycan strands from murein of <italic>Escherichia coli</italic> by reversed-phase high-performance liquid chromatography</article-title><source>Analytical Biochemistry</source><volume>190</volume><fpage>120</fpage><lpage>128</lpage><pub-id pub-id-type="doi">10.1016/0003-2697(90)90144-x</pub-id><pub-id pub-id-type="pmid">2285138</pub-id></element-citation></ref><ref id="bib14"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Juan</surname><given-names>C</given-names></name><name><surname>Torrens</surname><given-names>G</given-names></name><name><surname>Barceló</surname><given-names>IM</given-names></name><name><surname>Oliver</surname><given-names>A</given-names></name></person-group><year iso-8601-date="2018">2018</year><article-title>Interplay between peptidoglycan biology and virulence in gram-negative pathogens</article-title><source>Microbiology and Molecular Biology Reviews</source><volume>82</volume><elocation-id>e00033</elocation-id><pub-id pub-id-type="doi">10.1128/MMBR.00033-18</pub-id><pub-id pub-id-type="pmid">30209071</pub-id></element-citation></ref><ref id="bib15"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Kühner</surname><given-names>D</given-names></name><name><surname>Stahl</surname><given-names>M</given-names></name><name><surname>Demircioglu</surname><given-names>DD</given-names></name><name><surname>Bertsche</surname><given-names>U</given-names></name></person-group><year iso-8601-date="2014">2014</year><article-title>From cells to muropeptide structures in 24 h: Peptidoglycan mapping by UPLC-MS</article-title><source>Scientific Reports</source><volume>4</volume><elocation-id>7494</elocation-id><pub-id pub-id-type="doi">10.1038/srep07494</pub-id><pub-id pub-id-type="pmid">25510564</pub-id></element-citation></ref><ref id="bib16"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Liu</surname><given-names>X</given-names></name><name><surname>Biboy</surname><given-names>J</given-names></name><name><surname>Consoli</surname><given-names>E</given-names></name><name><surname>Vollmer</surname><given-names>W</given-names></name><name><surname>den Blaauwen</surname><given-names>T</given-names></name></person-group><year iso-8601-date="2020">2020</year><article-title>Mrec and MRED balance the interaction between the elongasome proteins pbp2 and RODA</article-title><source>PLOS Genetics</source><volume>16</volume><elocation-id>e1009276</elocation-id><pub-id pub-id-type="doi">10.1371/journal.pgen.1009276</pub-id><pub-id pub-id-type="pmid">33370261</pub-id></element-citation></ref><ref id="bib17"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Mainardi</surname><given-names>JL</given-names></name><name><surname>Villet</surname><given-names>R</given-names></name><name><surname>Bugg</surname><given-names>TD</given-names></name><name><surname>Mayer</surname><given-names>C</given-names></name><name><surname>Arthur</surname><given-names>M</given-names></name></person-group><year iso-8601-date="2008">2008</year><article-title>Evolution of peptidoglycan biosynthesis under the selective pressure of antibiotics in gram-positive bacteria</article-title><source>FEMS Microbiology Reviews</source><volume>32</volume><fpage>386</fpage><lpage>408</lpage><pub-id pub-id-type="doi">10.1111/j.1574-6976.2007.00097.x</pub-id><pub-id pub-id-type="pmid">18266857</pub-id></element-citation></ref><ref id="bib18"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Martin</surname><given-names>SA</given-names></name><name><surname>Rosenthal</surname><given-names>RS</given-names></name><name><surname>Biemann</surname><given-names>K</given-names></name></person-group><year iso-8601-date="1987">1987</year><article-title>Fast atom bombardment mass spectrometry and tandem mass spectrometry of biologically active peptidoglycan monomers from <italic>Neisseria gonorrhoeae</italic></article-title><source>The Journal of Biological Chemistry</source><volume>262</volume><fpage>7514</fpage><lpage>7522</lpage></element-citation></ref><ref id="bib19"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Morè</surname><given-names>N</given-names></name><name><surname>Martorana</surname><given-names>AM</given-names></name><name><surname>Biboy</surname><given-names>J</given-names></name><name><surname>Otten</surname><given-names>C</given-names></name><name><surname>Winkle</surname><given-names>M</given-names></name><name><surname>Serrano</surname><given-names>CKG</given-names></name><name><surname>Montón Silva</surname><given-names>A</given-names></name><name><surname>Atkinson</surname><given-names>L</given-names></name><name><surname>Yau</surname><given-names>H</given-names></name><name><surname>Breukink</surname><given-names>E</given-names></name><name><surname>den Blaauwen</surname><given-names>T</given-names></name><name><surname>Vollmer</surname><given-names>W</given-names></name><name><surname>Polissi</surname><given-names>A</given-names></name></person-group><year iso-8601-date="2019">2019</year><article-title>Peptidoglycan remodeling enables <italic>Escherichia coli</italic> to survive severe outer membrane assembly defect</article-title><source>MBio</source><volume>10</volume><elocation-id>e02729</elocation-id><pub-id pub-id-type="doi">10.1128/mBio.02729-18</pub-id><pub-id pub-id-type="pmid">30723128</pub-id></element-citation></ref><ref id="bib20"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Mudd</surname><given-names>S</given-names></name><name><surname>Lackman</surname><given-names>DB</given-names></name></person-group><year iso-8601-date="1941">1941</year><article-title>Bacterial morphology as shown by the electron microscope: I. Structural differentiation within the streptococcal cell</article-title><source>Journal of Bacteriology</source><volume>41</volume><fpage>415</fpage><lpage>420</lpage><pub-id pub-id-type="doi">10.1128/jb.41.3.415-420.1941</pub-id><pub-id pub-id-type="pmid">16560410</pub-id></element-citation></ref><ref id="bib21"><element-citation publication-type="software"><person-group person-group-type="author"><name><surname>Patel</surname><given-names>AV</given-names></name></person-group><year iso-8601-date="2021">2021</year><data-title>PGFinder</data-title><source>GitHub</source><ext-link ext-link-type="uri" xlink:href="https://github.com/Mesnage-Org/PGFinder">https://github.com/Mesnage-Org/PGFinder</ext-link></element-citation></ref><ref id="bib22"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Rogers</surname><given-names>HJ</given-names></name><name><surname>Perkins</surname><given-names>HR</given-names></name></person-group><year iso-8601-date="1959">1959</year><article-title>Cell-wall mucopeptides of <italic>Staphyloccus aureus</italic> and <italic>Micrococcus lysodeikticus</italic></article-title><source>Nature</source><volume>184</volume><fpage>520</fpage><lpage>524</lpage><pub-id pub-id-type="doi">10.1038/184520a0</pub-id><pub-id pub-id-type="pmid">14438367</pub-id></element-citation></ref><ref id="bib23"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Schleifer</surname><given-names>KH</given-names></name><name><surname>Kandler</surname><given-names>O</given-names></name></person-group><year iso-8601-date="1972">1972</year><article-title>Peptidoglycan types of bacterial cell walls and their taxonomic implications</article-title><source>Bacteriological Reviews</source><volume>36</volume><fpage>407</fpage><lpage>477</lpage><pub-id pub-id-type="doi">10.1128/br.36.4.407-477.1972</pub-id><pub-id pub-id-type="pmid">4568761</pub-id></element-citation></ref><ref id="bib24"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Tipper</surname><given-names>DJ</given-names></name></person-group><year iso-8601-date="2002">2002</year><article-title>Alkali-catalyzed elimination of D-lactic acid from muramic acid and its derivatives and the determination of muramic acid</article-title><source>Biochemistry</source><volume>7</volume><fpage>1441</fpage><lpage>1449</lpage><pub-id pub-id-type="doi">10.1021/bi00844a029</pub-id></element-citation></ref><ref id="bib25"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Turner</surname><given-names>RD</given-names></name><name><surname>Mesnage</surname><given-names>S</given-names></name><name><surname>Hobbs</surname><given-names>JK</given-names></name><name><surname>Foster</surname><given-names>SJ</given-names></name></person-group><year iso-8601-date="2018">2018</year><article-title>Molecular imaging of glycan chains couples cell-wall polysaccharide architecture to bacterial cell morphology</article-title><source>Nature Communications</source><volume>9</volume><elocation-id>1263</elocation-id><pub-id pub-id-type="doi">10.1038/s41467-018-03551-y</pub-id><pub-id pub-id-type="pmid">29593214</pub-id></element-citation></ref><ref id="bib26"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Tyanova</surname><given-names>S</given-names></name><name><surname>Temu</surname><given-names>T</given-names></name><name><surname>Sinitcyn</surname><given-names>P</given-names></name><name><surname>Carlson</surname><given-names>A</given-names></name><name><surname>Hein</surname><given-names>MY</given-names></name><name><surname>Geiger</surname><given-names>T</given-names></name><name><surname>Mann</surname><given-names>M</given-names></name><name><surname>Cox</surname><given-names>J</given-names></name></person-group><year iso-8601-date="2016">2016</year><article-title>The Perseus computational platform for comprehensive analysis of (prote)omics data</article-title><source>Nature Methods</source><volume>13</volume><fpage>731</fpage><lpage>740</lpage><pub-id pub-id-type="doi">10.1038/nmeth.3901</pub-id><pub-id pub-id-type="pmid">27348712</pub-id></element-citation></ref><ref id="bib27"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Vollmer</surname><given-names>W</given-names></name><name><surname>Blanot</surname><given-names>D</given-names></name><name><surname>de Pedro</surname><given-names>MA</given-names></name></person-group><year iso-8601-date="2008">2008</year><article-title>Peptidoglycan structure and architecture</article-title><source>FEMS Microbiology Reviews</source><volume>32</volume><fpage>149</fpage><lpage>167</lpage><pub-id pub-id-type="doi">10.1111/j.1574-6976.2007.00094.x</pub-id><pub-id pub-id-type="pmid">18194336</pub-id></element-citation></ref><ref id="bib28"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Watanabe</surname><given-names>Y</given-names></name><name><surname>Aoki-Kinoshita</surname><given-names>KF</given-names></name><name><surname>Ishihama</surname><given-names>Y</given-names></name><name><surname>Okuda</surname><given-names>S</given-names></name></person-group><year iso-8601-date="2021">2021</year><article-title>Glycopost realizes fair principles for glycomics mass spectrometry data</article-title><source>Nucleic Acids Research</source><volume>49</volume><fpage>D1523</fpage><lpage>D1528</lpage><pub-id pub-id-type="doi">10.1093/nar/gkaa1012</pub-id><pub-id pub-id-type="pmid">33174597</pub-id></element-citation></ref><ref id="bib29"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Weidel</surname><given-names>W</given-names></name><name><surname>Pelzer</surname><given-names>H</given-names></name></person-group><year iso-8601-date="1964">1964</year><article-title>Bagshaped macromolecules - a new outlook on bacterial cell walls</article-title><source>Advances in Enzymology and Related Subjects of Biochemistry</source><volume>26</volume><fpage>193</fpage><lpage>232</lpage><pub-id pub-id-type="doi">10.1002/9780470122716.ch5</pub-id></element-citation></ref><ref id="bib30"><element-citation publication-type="journal"><person-group person-group-type="author"><name><surname>Wheeler</surname><given-names>R</given-names></name><name><surname>Chevalier</surname><given-names>G</given-names></name><name><surname>Eberl</surname><given-names>G</given-names></name><name><surname>Gomperts Boneca</surname><given-names>I</given-names></name></person-group><year iso-8601-date="2014">2014</year><article-title>The biology of bacterial peptidoglycans and their impact on host immunity and physiology</article-title><source>Cellular Microbiology</source><volume>16</volume><fpage>1014</fpage><lpage>1023</lpage><pub-id pub-id-type="doi">10.1111/cmi.12304</pub-id><pub-id pub-id-type="pmid">24779390</pub-id></element-citation></ref></ref-list><app-group><app id="appendix-1"><title>Appendix 1</title><fig id="app1fig1" position="float"><label>Appendix 1—figure 1.</label><caption><title>UHPLC-MS chromatogram of <italic>E. coli</italic> reduced disaccharide peptides.</title></caption><graphic mime-subtype="tiff" mimetype="image" xlink:href="elife-70597-app1-fig1-v1.tif"/></fig><fig id="app1fig2" position="float"><label>Appendix 1—figure 2.</label><caption><title>Consistency of <italic>E. coli</italic> PG analyses.</title><p>(<bold>a</bold>) Pearson’s correlation coefficients across biological replicates of <italic>E. coli</italic> BW25113. (<bold>b</bold>) Muropeptide distribution according to degree of crosslinking. The crosslinking index was calculated as described previously (<xref ref-type="bibr" rid="bib11">Glauner, 1988</xref>). (<bold>c</bold>) Pairwise comparisons of intensities corresponding to individual muropeptides identified in biological replicates. WT1, WT2 and WT3 correspond to individual biological replicates; Av., average abundance; SD, standard deviation.</p></caption><graphic mime-subtype="tiff" mimetype="image" xlink:href="elife-70597-app1-fig2-v1.tif"/></fig></app><app id="appendix-2"><title>Appendix 2</title><fig id="app2fig1" position="float"><label>Appendix 2—figure 1.</label><caption><title>Workflow for production of MaxQuant compatible MS data files from Agilent QTOF data.</title><p>Agilent MS data (data: .d) is converted by Proteowizard to a mzML format (data: XML). Relevant settings for Proteowizard are shown (left). mzML file is then converted by TOPPAS to a mzXML file (data: XML). Relevant settings are shown (right).</p></caption><graphic mime-subtype="tiff" mimetype="image" xlink:href="elife-70597-app2-fig1-v1.tif"/></fig><fig id="app2fig2" position="float"><label>Appendix 2—figure 2.</label><caption><title>Workflow for MS data processing using MaxQuant, before automated analysis.</title><p>mzXML (data: XML) is passed to MaxQuant (process) for deconvolution and monoisotopic mass determination. Default values used except where indicated (right). MaxQuant output (data: text file) is then passed to the data parser module (process). This module removes superfluous data and reformats remaining data to be compatible with the matching script as an Excel file (data: xlsx).</p></caption><graphic mime-subtype="tiff" mimetype="image" xlink:href="elife-70597-app2-fig2-v1.tif"/></fig></app></app-group></back><sub-article article-type="decision-letter" id="sa1"><front-stub><article-id pub-id-type="doi">10.7554/eLife.70597.sa1</article-id><title-group><article-title>Decision letter</article-title></title-group><contrib-group content-type="section"><contrib contrib-type="editor"><name><surname>Blokesch</surname><given-names>Melanie</given-names></name><role>Reviewing Editor</role><aff><institution>Ecole Polytechnique Fédérale de Lausanne</institution><country>Switzerland</country></aff></contrib></contrib-group><contrib-group><contrib contrib-type="reviewer"><name><surname>Chou</surname><given-names>Seemay</given-names></name><role>Reviewer</role><aff><institution>UCSF</institution><country>United States</country></aff></contrib><contrib contrib-type="reviewer"><name><surname>Clarke</surname><given-names>Anthony</given-names></name><role>Reviewer</role></contrib></contrib-group></front-stub><body><boxed-text id="box1"><p>Our editorial process produces two outputs: i) <ext-link ext-link-type="uri" xlink:href="https://sciety.org/articles/activity/10.1101/2021.06.01.446515">public reviews</ext-link> designed to be posted alongside <ext-link ext-link-type="uri" xlink:href="https://www.biorxiv.org/content/10.1101/2021.06.01.446515v1">the preprint</ext-link> for the benefit of readers; ii) feedback on the manuscript for the authors, including requests for revisions, shown below. We also include an acceptance summary that explains what the editors found interesting or important about the work.</p></boxed-text><p><bold>Acceptance summary:</bold></p><p>This manuscript presents the development and validation of a new tool for the characterization of peptidoglycan (PG), the essential cell wall polymer of bacteria. PG is a single large macromolecule that protects almost all bacterial cells. The newly developed open access tool will greatly facilitate comparative quantitative analyses and the determination of compositional diversity of PG, which might ultimately contribute to the development of new antibacterials that target this essential cell wall component.</p><p><bold>Decision letter after peer review:</bold></p><p>Thank you for submitting your article &quot;PGfinder, a novel analysis pipeline for the consistent, reproducible and high-resolution structural analysis of bacterial peptidoglycans&quot; for consideration by <italic>eLife</italic>. Your article has been reviewed by 3 peer reviewers, and the evaluation has been overseen by a Reviewing Editor and Gisela Storz as the Senior Editor. The following individuals involved in review of your submission have agreed to reveal their identity: Seemay Chou (Reviewer #2); Anthony Clarke (Reviewer #3).</p><p>The reviewers have discussed their reviews with one another, and the Reviewing Editor has drafted this to help you prepare a revised submission.</p><p>Essential revisions:</p><p>The authors are encouraged to revise their manuscript by properly stating two weaknesses of their tool (as it stands in its current form):</p><p>1) The lack of important PG modifications in the mass spectrometry-based library of the muropeptides such as O-acetylation; inclusion of such modification would be required for the broader application of the analysis tool (e.g., as those modifications are frequently encountered in pathogenic Gram-positive bacteria);</p><p>2) The emphasis on comparative analyses between different samples (e.g., from the same species/strain but grown under different conditions) versus absolute quantification using MS, which might be more challenging. Future users are therefore encouraged to include appropriate controls.</p><p><italic>Reviewer #1 (Recommendations for the authors):</italic></p><p>Lines 285-288 say, &quot;The present study confirmed that the complexity of bacterial PGs was greatly underestimated until now. Our unbiased search identified &gt;106 masses matching <italic>E. coli</italic> muropeptide structures representing a number strikingly larger than previously reported (Kuhner, et al., 2014).&quot; This may be an overestimate of the increase in complexity revealed by the new method of analysis. While older methods certainly identified fewer muropeptides, I think it was understood that the range of modifications and crosslinked dimers and trimers was greater than what was identified, and that many low-abundance muropeptides were present. Many of these could be predicted. This does not diminish the importance of this tool for allowing the rapid identification and quantification of many of these assumed muropeptides, but the actual increased demonstration of complexity should not be overstated.</p><p>Lines 298-301: The authors describe the advantages of muropeptide quantification using XIC, such as resolving overlapping peaks. There should also be a discussion of challenges of quantitative analysis using MS data, such as potential differences in ionization efficiency between muropeptides with significantly different peptide side chains, with possible lower detection of larger muropeptides due to lessor ionization, or differences of in-source fragmentation. Older methods using UV absorbance quantification were likely more accurate for quantification across the diversity of identified muropeptides, but missed low abundance muropeptides.</p><p>Lines 305-306: I agree that this represents a powerful tool for qualitative and reproducible PG analyses. The demonstration of absolute quantification is not clear. Yes, quantitative differences between the C. difficile strains is very clear, and the reproducibility of differences between the measured <italic>E. coli</italic> average cross-linking between this method and published values is clear, but whether the lower cross-linking determined by this method are more accurate is not clear. Is this due to greater detection of minor muropeptides or more accurate quantification of all muropeptides species? Or is the difference due to lessor detection of ions for the cross-linked muropeptides?</p><p>Lines 369-370 say, &quot;Samples were diluted to contain 150mAU/μl of the major monomer. Based on the dry weight of the PG sample, we estimated that this corresponded to approximately 50μg of material&quot; It is not clear what this 50 µg is. Does this dilution produce a suspension of 50 µg/µl? Does the resulting 10 µl injected for UHPLC contain 50 µg?</p><p><italic>Reviewer #2 (Recommendations for the authors):</italic></p><p>Below are several comments/questions on the paper as a whole but also on some specific sections.</p><p>1. The workflow was explained in detail in the text and also through the use of helpful schematics</p><p>2. Automating the analysis of MS cell wall data is a major advance in the field, especially through the use of open-source software.</p><p>3. Line 147, do you think the remaining 52-59% of total ion intensity is noise? If so, why so much? Or do you think there is some real signal that you are missing? If so, what could be improved to capture that?</p><p>4. Lines 149, what stereoisomers are you referring to? How do they form? Do you think these form in living cells?</p><p>5. Lines 158, 165, and 169, the authors found some discrepancies between current knowledge on <italic>E. coli</italic> PG and their data.</p><p>a. They stated in the text that further biochemical and molecular biology studies would be required to validate their findings.</p><p>6. Line 236, why did you increase the mass tolerance here to 25ppm? What was the problem if you used 10ppm as you did for the previous analysis?</p><p>7. The authors established the application of an alternative open-source deconvolution software.</p><p>a. They showed that it works as well as the commercial counterparts.</p><p>b. Lines 244-246, they described variability issues with the deconvolution step – what could be done to improve the software used here? This could help with the application of your approach across more datasets.</p><p>8. The authors provided multiple schematics that are helpful to visualize their workflow.</p><p>9. Their previous study and this one demonstrate the capacity for discovery when this type of unbiased approach was applied.</p><p>a. Especially regarding their novel findings for <italic>E. coli</italic> and C. difficile PG composition.</p><p>10. Their approach has advantages such as open-access and availability as a Jupyter notebook that would make it broadly available in the community. However, they also describe issues that need to be addressed.</p><p>a. One such issue is the ability to analyze MS/MS data.</p><p>11. The paper establishes high standards for other studies relying on MS data but also in general for other cell wall papers relying on chromatographic techniques</p><p>a. Biological replicates, list of searched muropeptides, and general data reporting</p><p>12. Some of the methods rely too heavily on previously published protocols. Could you provide brief details on the steps you followed in this study?</p><p>a. Lines 354, 361, and 395.</p><p>13. Check spellings of meso-diaminopimelic acid – it varies across the paper.</p><p>14. Correct g-D-Glutamate to γ-D-Glutamate – several instances across the paper.</p><p><italic>Reviewer #3 (Recommendations for the authors):</italic></p><p>1. Lines 125-126: Why did the authors limit the modifications to PG in their Library 3 to only those listed on these lines. Importantly, given that all pathogenic Gram-positive bacteria, and many Gram-negative bacteria produce O-acetylated PG, this reviewer is surprised that this modification was not (apparently) included, especially since reference is made to it on line included in Line 332 of the Discussion. This addition would be necessary for the analysis of the PG from these bacteria given the levels of O-acetylation which would greatly influence % compositional analyses of the various muropeptides.</p><p>2. Throughout the manuscript, including main and supplemental figures and tables, the authors have used J as their one-letter code abbreviation for diaminopimelic acid. Unfortunately, the IUPAC-IUB Joint Commission on Biochemical Nomenclature define J as the one-letter code for the combination of either isoleucine or leucine (I/L) ; please see, e.g., <ext-link ext-link-type="uri" xlink:href="http://www.insdc.org/documents/feature_table.html#7.4.3">http://www.insdc.org/documents/feature_table.html#7.4.3</ext-link>.</p><p>All letters of the Latin alphabet are already used, and the Joint Commission states Dpm as the abbreviation for diaminopimelic acid. Using J as the (defined) abbreviation may have been acceptable with its isolated use, but it is very problematic when combined with the conventional code for the other amino acids. Clearly, using the three-letter abbreviation in combination with the one-letter code is awkward (particularly when convention usually requires the use of one or the other, not their combination) but in this instance, maybe the only option. Unless, the authors want to consider introducing another unconventional approach and move to eg., Greek and use delta for Dpm. This maybe going too far, but I can't think of any other alternative.</p></body></sub-article><sub-article article-type="reply" id="sa2"><front-stub><article-id pub-id-type="doi">10.7554/eLife.70597.sa2</article-id><title-group><article-title>Author response</article-title></title-group></front-stub><body><disp-quote content-type="editor-comment"><p>Essential revisions:</p><p>The authors are encouraged to revise their manuscript by properly stating two weaknesses of their tool (as it stands in its current form):</p><p>1) The lack of important PG modifications in the mass spectrometry-based library of the muropeptides such as O-acetylation; inclusion of such modification would be required for the broader application of the analysis tool (e.g., as those modifications are frequently encountered in pathogenic Gram-positive bacteria);</p></disp-quote><p>We agree that this was a limitation of the version of PGFinder we made available with the original submission. Rather than stating this weakness in the revised manuscript, we have added this functionality to PGFinder (now v.0.02; https://github.com/Mesnage-Org/PGFinder/releases/tag/v0.02). Users now have the possibility to search for O-acetylated PG fragments.</p><p>It is important to point out that given the complexity of PG structure, PGFinder will continue to be developed. We are very keen to receive feedback from users; the Jupyter notebook allows them to do so (see <xref ref-type="fig" rid="sa2fig1">Author response image 1</xref>).</p><fig id="sa2fig1" position="float"><label>Author response image 1.</label><graphic mime-subtype="tiff" mimetype="image" xlink:href="elife-70597-sa2-fig1-v1.tif"/></fig><disp-quote content-type="editor-comment"><p>2) The emphasis on comparative analyses between different samples (e.g., from the same species/strain but grown under different conditions) versus absolute quantification using MS, which might be more challenging. Future users are therefore encouraged to include appropriate controls.</p></disp-quote><p>We agree with this point, but we suggest that relative quantification of muropeptides is sufficient for most experimental designs that researchers might wish to use. The quantification of, for example, proteomes or metabolomes by mass spectrometry is rarely absolute as relative quantification is sufficient to measure differences in the abundance of analytes across a wide range of sample types. Relative quantification, when performed with the appropriate number of replicates and robust statistical analysis, allows for confident measurement of changes in the abundance of muropeptides.</p><p>We have clearly stated that our data analysis pipeline is compatible with relative quantification in the revised manuscript (L. 304-307) and put emphasis on the fact that PGFinder is an ideal tool for comparative analyses, the most common type of PG analyses carried out to date.</p><disp-quote content-type="editor-comment"><p>Reviewer #1 (Recommendations for the authors):</p><p>Lines 285-288 say, &quot;The present study confirmed that the complexity of bacterial PGs was greatly underestimated until now. Our unbiased search identified &gt;106 masses matching <italic>E. coli</italic> muropeptide structures representing a number strikingly larger than previously reported (Kuhner, et al., 2014).&quot; This may be an overestimate of the increase in complexity revealed by the new method of analysis. While older methods certainly identified fewer muropeptides, I think it was understood that the range of modifications and crosslinked dimers and trimers was greater than what was identified, and that many low-abundance muropeptides were present. Many of these could be predicted. This does not diminish the importance of this tool for allowing the rapid identification and quantification of many of these assumed muropeptides, but the actual increased demonstration of complexity should not be overstated.</p></disp-quote><p>We agree with the reviewer’s comment. The point we tried to make was that PGFinder allows this complexity to be revealed, not that we discovered something conceptually novel. We have clearly stated this in the revised manuscript that the complexity described here was expected (L. 284-288).</p><disp-quote content-type="editor-comment"><p>Lines 298-301: The authors describe the advantages of muropeptide quantification using XIC, such as resolving overlapping peaks. There should also be a discussion of challenges of quantitative analysis using MS data, such as potential differences in ionization efficiency between muropeptides with significantly different peptide side chains, with possible lower detection of larger muropeptides due to lessor ionization, or differences of in-source fragmentation. Older methods using UV absorbance quantification were likely more accurate for quantification across the diversity of identified muropeptides, but missed low abundance muropeptides.</p></disp-quote><p>We discussed the limitations of our approach for quantitative analysis but insisted on the fact that the primary application of PGFinder is to carry out comparative analyses (L. 304-307). As we perform relative quantification of individual muropeptides across different samples, the relative ionisation of different muropeptides relative to each other is irrelevant. In-source fragmentation is handled by PGFinder and taken into account to consolidate intensities during the “clean up step”.</p><disp-quote content-type="editor-comment"><p>Lines 305-306: I agree that this represents a powerful tool for qualitative and reproducible PG analyses. The demonstration of absolute quantification is not clear. Yes, quantitative differences between the C. difficile strains is very clear, and the reproducibility of differences between the measured <italic>E. coli</italic> average cross-linking between this method and published values is clear, but whether the lower cross-linking determined by this method are more accurate is not clear. Is this due to greater detection of minor muropeptides or more accurate quantification of all muropeptides species? Or is the difference due to lessor detection of ions for the cross-linked muropeptides?</p></disp-quote><p>As stated previously, PGFinder is primarily a tool for comparative analyses.</p><p>The cross-linking index reported in this study is relatively low as compared to the one described in the paper cited. The seminal work of Glauner (1988) contains little information about the methods used to ascertain muropeptides identity. The interpretation of such discrepancies therefore seems highly speculative, and we would rather not try to address this point in the revised manuscript to avoid confusing the reader.</p><p>We feel that the issues with XIC quantification (highlighted in the revised version) could account for these differences.</p><p>Lines 369-370 say, &quot;Samples were diluted to contain 150mAU/μl of the major monomer. Based on the dry weight of the PG sample, we estimated that this corresponded to approximately 50μg of material&quot; It is not clear what this 50 µg is. Does this dilution produce a suspension of 50 µg/µl? Does the resulting 10 µl injected for UHPLC contain 50 µg?50µg of material was analysed by UHPLC-MS (L. 378-379).</p><disp-quote content-type="editor-comment"><p>Reviewer #2 (Recommendations for the authors):</p><p>Below are several comments/questions on the paper as a whole but also on some specific sections.</p><p>1. The workflow was explained in detail in the text and also through the use of helpful schematics.</p><p>2. Automating the analysis of MS cell wall data is a major advance in the field, especially through the use of open-source software.</p></disp-quote><p>No comments required.</p><disp-quote content-type="editor-comment"><p>3. Line 147, do you think the remaining 52-59% of total ion intensity is noise? If so, why so much? Or do you think there is some real signal that you are missing? If so, what could be improved to capture that?</p></disp-quote><p>The PG extraction protocol we use does not involve a DNAse/RNase treatment or a step to break the cells. Our prior AFM experiments on sacculi prepared with this method clearly show that some material can be trapped inside the PG sacculi (doi:10.1111/j.1365-2958.2011.07871; DOI: 10.1038/ncomms2503). The enzyme used for the final digestion step is also present in the mixture.</p><disp-quote content-type="editor-comment"><p>4. Lines 149, what stereoisomers are you referring to? How do they form? Do you think these form in living cells?</p></disp-quote><p>The presence of stereoisomers has been reported in the literature (DOI: 10.1038/srep07494). These correspond to muropeptides that contain a small proportion of isomers that are incorporated in peptidoglycan peptide stems. Stereoisomers are therefore formed in living cells and are not an artefact associated with the chromatographic techniques used.</p><disp-quote content-type="editor-comment"><p>5. Lines 158, 165, and 169, the authors found some discrepancies between current knowledge on <italic>E. coli</italic> PG and their data.</p><p>a. They stated in the text that further biochemical and molecular biology studies would be required to validate their findings.</p></disp-quote><p>No comment required.</p><disp-quote content-type="editor-comment"><p>6. Line 236, why did you increase the mass tolerance here to 25ppm? What was the problem if you used 10ppm as you did for the previous analysis?</p></disp-quote><p>We set the mass tolerance at 25 ppm to use the same threshold as described by Anderson et al. This comment has been added in the revised manuscript (L.235).</p><disp-quote content-type="editor-comment"><p>7. The authors established the application of an alternative open-source deconvolution software.</p><p>a. They showed that it works as well as the commercial counterparts.</p><p>b. Lines 244-246, they described variability issues with the deconvolution step – what could be done to improve the software used here? This could help with the application of your approach across more datasets</p></disp-quote><p>The two deconvolution software rely on distinct methods; Byos uses a machine learning algorithm to identify an isotopic series and then determine the charge state and m0 ion whereas MaxQuant uses a system based on constructing a graph network of the <italic>m/z</italic> spectra and identifying disconnected clusters in these networks to identify an isotopic series before determining charge state and m0 masses.</p><p>Neither software is open-source so we cannot improve them.</p><disp-quote content-type="editor-comment"><p>8. The authors provided multiple schematics that are helpful to visualize their workflow.</p><p>9. Their previous study and this one demonstrate the capacity for discovery when this type of unbiased approach was applied.</p><p>a. Especially regarding their novel findings for <italic>E. coli</italic> and C. difficile PG composition.</p></disp-quote><p>No comments required.</p><disp-quote content-type="editor-comment"><p>10. Their approach has advantages such as open-access and availability as a Jupyter notebook that would make it broadly available in the community. However, they also describe issues that need to be addressed.</p><p>a. One such issue is the ability to analyze MS/MS data.</p></disp-quote><p>This point is discussed in the manuscript (L.330-333).</p><disp-quote content-type="editor-comment"><p>11. The paper establishes high standards for other studies relying on MS data but also in general for other cell wall papers relying on chromatographic techniques</p><p>a. Biological replicates, list of searched muropeptides, and general data reporting</p></disp-quote><p>No comment required.</p><disp-quote content-type="editor-comment"><p>12. Some of the methods rely too heavily on previously published protocols. Could you provide brief details on the steps you followed in this study?</p><p>a. Lines 354, 361, and 395.</p></disp-quote><p>Reduction and β-elimination are briefly described (L. 364-370).</p><p>PG desalting is described in great details so no more details can be added!</p><p>The analysis using Perseus is widely used by proteomics users and all relevant details are provided. Many tutorials are available to explain the entire process in great detail.</p><disp-quote content-type="editor-comment"><p>13. Check spellings of meso-diaminopimelic acid – it varies across the paper.</p></disp-quote><p>Done.</p><disp-quote content-type="editor-comment"><p>14. Correct g-D-Glutamate to γ-D-Glutamate – several instances across the paper.</p></disp-quote><p>Done.</p><disp-quote content-type="editor-comment"><p>Reviewer #3 (Recommendations for the authors):</p><p>1. Lines 125-126: Why did the authors limit the modifications to PG in their Library 3 to only those listed on these lines. Importantly, given that all pathogenic Gram-positive bacteria, and many Gram-negative bacteria produce O-acetylated PG, this reviewer is surprised that this modification was not (apparently) included, especially since reference is made to it on line included in Line 332 of the Discussion. This addition would be necessary for the analysis of the PG from these bacteria given the levels of O-acetylation which would greatly influence % compositional analyses of the various muropeptides.</p></disp-quote><p>We have modified PGFinder to include this modification as a search option.</p><disp-quote content-type="editor-comment"><p>2. Throughout the manuscript, including main and supplemental figures and tables, the authors have used J as their one-letter code abbreviation for diaminopimelic acid. Unfortunately, the IUPAC-IUB Joint Commission on Biochemical Nomenclature define J as the one-letter code for the combination of either isoleucine or leucine (I/L) ; please see, e.g., <ext-link ext-link-type="uri" xlink:href="http://www.insdc.org/documents/feature_table.html#7.4.3">http://www.insdc.org/documents/feature_table.html#7.4.3</ext-link>.</p><p>All letters of the Latin alphabet are already used, and the Joint Commission states Dpm as the abbreviation for diaminopimelic acid. Using J as the (defined) abbreviation may have been acceptable with its isolated use, but it is very problematic when combined with the conventional code for the other amino acids. Clearly, using the three-letter abbreviation in combination with the one-letter code is awkward (particularly when convention usually requires the use of one or the other, not their combination) but in this instance, maybe the only option. Unless, the authors want to consider introducing another unconventional approach and move to eg., Greek and use δ for Dpm. This maybe going too far, but I can't think of any other alternative.</p></disp-quote><p>Unfortunately, we cannot use Greek letters because the.csv files used by PGFinder (the muropeptide databases) do not support symbols or special characters. We argue that this is not a problem since the final output (the muropeptide table describing PG structure) can always be modified using the “search/replace” function. The letter J can therefore be replaced at the final stages of the PG analysis by any letter(s)/symbol.</p><p>To be consistent with the previous work (published by the reviewer) we have replaced the letter “J” by “m” in Table 1 (<italic>E. coli</italic> muropeptide table), Table 2 (<italic>P. aeruginosa</italic> muropeptide table), Figure 3 (Sankey diagram) and Figure 4 (<italic>C. difficile</italic> volcano plot). We added a sentence in the manuscript to explain this (L. 143-146).</p></body></sub-article></article>