{\rtf1\ansi\ansicpg1252\uc1 \deff0\deflang1033\deflangfe1033{\fonttbl{\f0\froman\fcharset0\fprq2{\*\panose 02020603050405020304}Times New Roman;}{\f2\fmodern\fcharset0\fprq1{\*\panose 02070309020205020404}Courier New;}
{\f45\fnil\fcharset2\fprq2{\*\panose 00000000000000000000}MS Reference 2;}{\f46\froman\fcharset238\fprq2 Times New Roman CE;}{\f47\froman\fcharset204\fprq2 Times New Roman Cyr;}{\f49\froman\fcharset161\fprq2 Times New Roman Greek;}
{\f50\froman\fcharset162\fprq2 Times New Roman Tur;}{\f51\froman\fcharset186\fprq2 Times New Roman Baltic;}{\f58\fmodern\fcharset238\fprq1 Courier New CE;}{\f59\fmodern\fcharset204\fprq1 Courier New Cyr;}{\f61\fmodern\fcharset161\fprq1 Courier New Greek;}
{\f62\fmodern\fcharset162\fprq1 Courier New Tur;}{\f63\fmodern\fcharset186\fprq1 Courier New Baltic;}}{\colortbl;\red0\green0\blue0;\red0\green0\blue255;\red0\green255\blue255;\red0\green255\blue0;\red255\green0\blue255;\red255\green0\blue0;
\red255\green255\blue0;\red255\green255\blue255;\red0\green0\blue128;\red0\green128\blue128;\red0\green128\blue0;\red128\green0\blue128;\red128\green0\blue0;\red128\green128\blue0;\red128\green128\blue128;\red192\green192\blue192;}{\stylesheet{
\widctlpar\adjustright \lang2057\cgrid \snext0 Normal;}{\s6\keepn\widctlpar\adjustright \i\lang2057\cgrid \sbasedon0 \snext0 heading 6;}{\*\cs10 \additive Default Paragraph Font;}{\s15\widctlpar\adjustright \f2\fs20\lang2057\cgrid \sbasedon0 \snext15 
Plain Text;}{\s16\widctlpar\adjustright \fs20\lang2057\cgrid \sbasedon0 \snext16 footnote text;}{\*\cs17 \additive \super \sbasedon10 footnote reference;}{\s18\widctlpar\adjustright \fs20\lang2057\cgrid \sbasedon0 \snext18 endnote text;}{
\s19\widctlpar\adjustright \i\lang2057\cgrid \sbasedon0 \snext19 Body Text;}}{\*\listtable{\list\listtemplateid134807567\listsimple{\listlevel\levelnfc0\leveljc0\levelfollow0\levelstartat1\levelspace0\levelindent0{\leveltext\'02\'00.;}{\levelnumbers\'01;}
\fbias0 \fi-360\li360\jclisttab\tx360 }{\listname ;}\listid43023407}{\list\listtemplateid-777621056\listsimple{\listlevel\levelnfc4\leveljc0\levelfollow0\levelstartat1\levelspace0\levelindent0{\leveltext\'02\'00.;}{\levelnumbers\'01;}\ulnone\fbias0 
\fi-360\li360\jclisttab\tx360 }{\listname ;}\listid426967849}{\list\listtemplateid134807567\listsimple{\listlevel\levelnfc0\leveljc0\levelfollow0\levelstartat1\levelspace0\levelindent0{\leveltext\'02\'00.;}{\levelnumbers\'01;}\fbias0 \fi-360\li360
\jclisttab\tx360 }{\listname ;}\listid587345553}{\list\listtemplateid-2050054196\listsimple{\listlevel\levelnfc4\leveljc0\levelfollow0\levelstartat1\levelspace0\levelindent0{\leveltext\'02\'00.;}{\levelnumbers\'01;}\b\fbias0 \fi-360\li360\jclisttab\tx360 
}{\listname ;}\listid637344700}{\list\listtemplateid134807567\listsimple{\listlevel\levelnfc0\leveljc0\levelfollow0\levelstartat1\levelspace0\levelindent0{\leveltext\'02\'00.;}{\levelnumbers\'01;}\fi-360\li360\jclisttab\tx360 }{\listname 
;}\listid1058282921}{\list\listtemplateid134807567\listsimple{\listlevel\levelnfc0\leveljc0\levelfollow0\levelstartat1\levelspace0\levelindent0{\leveltext\'02\'00.;}{\levelnumbers\'01;}\fbias0 \fi-360\li360\jclisttab\tx360 }{\listname ;}\listid1650401830}
{\list\listtemplateid134807567\listsimple{\listlevel\levelnfc0\leveljc0\levelfollow0\levelstartat1\levelspace0\levelindent0{\leveltext\'02\'00.;}{\levelnumbers\'01;}\fbias0 \fi-360\li360\jclisttab\tx360 }{\listname ;}\listid2051342558}}
{\*\listoverridetable{\listoverride\listid426967849\listoverridecount0\ls1}{\listoverride\listid637344700\listoverridecount0\ls2}{\listoverride\listid2051342558\listoverridecount0\ls3}{\listoverride\listid43023407\listoverridecount0\ls4}
{\listoverride\listid1650401830\listoverridecount0\ls5}{\listoverride\listid1058282921\listoverridecount0\ls6}{\listoverride\listid587345553\listoverridecount0\ls7}}{\info{\title Introducing COMPARA, the Portuguese-English parallel1 corpus}
{\author Frankenberg Garcia}{\operator Diana Santos}{\creatim\yr2001\mo2\dy26\hr16\min2}{\revtim\yr2001\mo2\dy26\hr16\min2}{\printim\yr2000\mo12\dy18\hr9\min49}{\version2}{\edmins0}{\nofpages3}{\nofwords3704}{\nofchars20002}{\*\company  }
{\nofcharsws25928}{\vern89}}\paperw11900\paperh16840\margl1418\margr1418\margt1418\margb1418 \widowctrl\ftnbj\aenddoc\makebackup\formshade\viewkind1\viewscale100\pgbrdrhead\pgbrdrfoot \fet0\sectd 
\psz9\sbknone\linex0\headery737\footery851\colsx709\endnhere\sectdefaultcl {\*\pnseclvl1\pnucrm\pnstart1\pnindent720\pnhang{\pntxta .}}{\*\pnseclvl2\pnucltr\pnstart1\pnindent720\pnhang{\pntxta .}}{\*\pnseclvl3\pndec\pnstart1\pnindent720\pnhang{\pntxta .}}
{\*\pnseclvl4\pnlcltr\pnstart1\pnindent720\pnhang{\pntxta )}}{\*\pnseclvl5\pndec\pnstart1\pnindent720\pnhang{\pntxtb (}{\pntxta )}}{\*\pnseclvl6\pnlcltr\pnstart1\pnindent720\pnhang{\pntxtb (}{\pntxta )}}{\*\pnseclvl7\pnlcrm\pnstart1\pnindent720\pnhang
{\pntxtb (}{\pntxta )}}{\*\pnseclvl8\pnlcltr\pnstart1\pnindent720\pnhang{\pntxtb (}{\pntxta )}}{\*\pnseclvl9\pnlcrm\pnstart1\pnindent720\pnhang{\pntxtb (}{\pntxta )}}\pard\plain \s15\widctlpar\outlinelevel0\adjustright \f2\fs20\lang2057\cgrid {\b\f0\fs28 
Introducing COMPARA, the Portuguese-English parallel}{\b\f0\fs28\super 1}{\b\outl\f0\fs28\super  }{\b\f0\fs28 corpus
\par }{\f0\fs24 Ana Frankenberg-Garcia (ISLA, Lisbon) & Diana Santos (SINTEF, Oslo)
\par }\pard\plain \widctlpar\adjustright \lang2057\cgrid {
\par }\pard\plain \s19\widctlpar\adjustright \i\lang2057\cgrid {\fs22 
\par }\pard \s19\widctlpar\adjustright {\fs22 This paper is an introduction to COMPARA. COMPARA is a machine-searchable, open-ended collection of Portuguese-English and Eng
lish-Portuguese source texts and translations. It was made for people who have never used corpora before as well as for experienced corpus users. COMPARA\rquote s encoding and alignment criteria allow users to inspect translators\rquote 
 notes and to investigate when an
d where translators have chosen to join, separate, delete, add and reorder sentences. Another innovative feature is that the corpus admits more than one translation per source text. COMPARA is encoded according to the IMS Corpus Workbench system and is fr
eely accessible on the WWW via the DISPARA interface. 
\par }\pard\plain \s15\widctlpar\adjustright \f2\fs20\lang2057\cgrid {\f0\fs24 
\par }\pard \s15\widctlpar\outlinelevel0\adjustright {\b\f0\fs24 
\par Basic characteristics of COMPARA
\par }\pard \s15\widctlpar\adjustright {\f0\fs24 This paper is an introduction to COMPARA, the Portuguese-English parallel corpus. Modelling itself on the part of the English-Norwegian Parallel Corpus devoted to par
allel texts (Johansson et al. 1999), COMPARA is a machine-readable and searchable collection of texts originally written in Portuguese and in English that have been aligned with their respective English and Portuguese translations. 
\par 
\par The basic characteristics of COMPARA are that it is:
\par 1. open-ended 
\par 2. for people who are not necessarily corpus-literate as well as for experienced corpus users
\par 3. searchable via the Internet 
\par 
\par 
\par The decision to leave COMPARA open-ended was taken partly so that it could gr
ow in whichever direction proved to become important to its users, and partly because this meant the texts incorporated in the corpus could be put to use as soon as they were processed. The second of these two reasons is not trivial: it meant that it was 
possible for the corpus to become operational within a reasonable amount of time. The trade-off, of course, is that at the time this paper was written (just over 
one year after the project began), COMPARA did not lend itself to analyses requiring large and representative language samples. 
\par 
\par COMPARA was made for anyone interested in the study of Portuguese and English contrasts. Potential users include Portuguese learners of English, English learners of Portuguese, students and teachers of translation, professi
onal translators, bilingual dictionary makers, developers of machine translation software and whoever else might be interested in translation language in and in the similarities and differences between Portuguese and English. 
\par 
\par One of our main concerns was to make sure that COMPARA could be used not only by experienced corpus users, but also by people who have never used a corpus before. 
\par 
\par Access to COMPARA is provided free of charge at: 
\par 
\par http://www.portugues.mct.pt/COMPARA/Welcome.html 
\par 
\par The above site is maintained by the Computational Processing of Portuguese project, which is also responsible for distributing other Portuguese language resources apart from COMPARA}{\f0\fs24\super 2}{\f0\fs24 . 
\par 
\par Every page in the COMPARA website is available in both Portuguese and English, so that people with very little Portuguese or very little English can still read them. 
\par 
\par }\pard \s15\widctlpar\outlinelevel0\adjustright {\b\f0\fs24 Text Selection
\par }\pard \s15\widctlpar\adjustright {\f0\fs24 When selecting texts for the corpus, all varieties of Portuguese and English were considered, and no priority was given to any particular variety. In terms of dat
e of publication, both contemporary and non-contemporary texts were accepted. In addition to this, the possibility of having a source text aligned with more than one translation was not ruled out. Having established this, it was decided to begin the corpu
s by assembling an initial collection of published fiction, although other genres are to be included in the corpus at a later stage. 
\par 
\par }\pard \s15\widctlpar\outlinelevel0\adjustright {\b\f0\fs24 Copyright permissions 
\par }\pard \s15\widctlpar\adjustright {\f0\fs24 For the initial, fiction part of the corpus, most efforts have so far been directed at obtaining co
pyright clearance in the Portuguese to English direction, because of the fewer Portuguese to English translations available to choose from. However, since COMPARA is to remain open-ended, obtaining a balanced corpus is not crucial.  The responsibility of 
achieving balance is in fact being deliberately transferred to the users of COMPARA, who are expected to pre-select the texts they want to use for their search queries if and when the question of balance is important.
\par 
\par Although not all copyright permissions
 applied to were granted, the overall response from authors, translators and publishers was quite encouraging, especially when considering the fact that they actually gave permission for their texts to be freely searchable on the Internet.
\par 
\par At the time this paper was written, COMPARA had permission to include extracts - usually 30% of a complete work}{\f0\fs24\super 3}{\f0\fs24 
 - of 60 different Portuguese-English text-pairs by authors and translators from Angola, Brazil, Mozambique, Portugal, South Africa, the United Kingdom and the
 United States. These texts represent the combined product of the work of 33 different authors and 31 translators}{\f0\fs24\super 4}{\f0\fs24 . 
\par 
\par Because COMPARA allows for the inclusion of more than one translation of the same source, some interesting text-pair combinations have eme
rged. For example, permission has been obtained to include extracts of a couple of novels by David Lodge (Lodge 1975; 1995a) paired up with both their Portuguese and Brazilian translations (respectively, Lodge 1995b, 1995c, 1997,
 1998), which can be useful
 for the study of similarities and differences between Brazilian and European Portuguese. Another interesting example is that of a Brazilian nineteenth century Romantic classic, Iracema (Alencar 1865), which has been paired up with a contemporary English 
translation published by Oxford University Press  only a few months ago }{\f0\fs24 (Alencar 2000)}{\f0\fs24  and a contemporaneous translation which dates back to 1886 (Alencar 1886) - this could be interesting for a diachronic study of translation.

\par 
\par }\pard \s15\widctlpar\outlinelevel0\adjustright {\b\f0\fs24 Corpus composition in December 2000
\par }\pard \s15\widctlpar\adjustright {\f0\fs24 The COMPARA corpus project began in mid-October 1999, and very few texts had been fully processed at the time this paper was written. The part of the corpus that was available for research in December
 2000 is summarized in table 1.
\par 
\par }{\f0\fs24\ul \page Table 1:Composition of COMPARA in December 2000}{\f0\fs24 
\par 
\par }\trowd \trqc\trgaph108\trleft-108\trbrdrt\brdrs\brdrw10 \trbrdrl\brdrs\brdrw10 \trbrdrb\brdrs\brdrw10 \trbrdrr\brdrs\brdrw10 \trbrdrh\brdrs\brdrw10 \trbrdrv\brdrs\brdrw10 \clvertalt\clbrdrt\brdrs\brdrw10 \clbrdrl\brdrs\brdrw10 \clbrdrb\brdrs\brdrw10 
\clbrdrr\brdrs\brdrw10 \cltxlrtb \cellx1843\clvertalt\clbrdrt\brdrs\brdrw10 \clbrdrl\brdrs\brdrw10 \clbrdrb\brdrs\brdrw10 \clbrdrr\brdrs\brdrw10 \cltxlrtb \cellx3402\clvertalt\clbrdrt\brdrs\brdrw10 \clbrdrl\brdrs\brdrw10 \clbrdrb\brdrs\brdrw10 \clbrdrr
\brdrs\brdrw10 \cltxlrtb \cellx5103\clvertalt\clbrdrt\brdrs\brdrw10 \clbrdrl\brdrs\brdrw10 \clbrdrb\brdrs\brdrw10 \clbrdrr\brdrs\brdrw10 \cltxlrtb \cellx6237\pard \s15\widctlpar\intbl\adjustright {\f0\fs24 COMPARA December 2000\cell }{\b\f0\fs24 
Portuguese language\cell English Language\cell Total\cell }\pard\plain \widctlpar\intbl\adjustright \lang2057\cgrid {\row }\trowd \trqc\trgaph108\trleft-108\trbrdrt\brdrs\brdrw10 \trbrdrl\brdrs\brdrw10 \trbrdrb\brdrs\brdrw10 \trbrdrr\brdrs\brdrw10 
\trbrdrh\brdrs\brdrw10 \trbrdrv\brdrs\brdrw10 \clvertalt\clbrdrt\brdrs\brdrw10 \clbrdrl\brdrs\brdrw10 \clbrdrb\brdrs\brdrw10 \clbrdrr\brdrs\brdrw10 \cltxlrtb \cellx1843\clvertalt\clbrdrt\brdrs\brdrw10 \clbrdrl\brdrs\brdrw10 \clbrdrb\brdrs\brdrw10 \clbrdrr
\brdrs\brdrw10 \cltxlrtb \cellx3402\clvertalt\clbrdrt\brdrs\brdrw10 \clbrdrl\brdrs\brdrw10 \clbrdrb\brdrs\brdrw10 \clbrdrr\brdrs\brdrw10 \cltxlrtb \cellx5103\clvertalt\clbrdrt\brdrs\brdrw10 \clbrdrl\brdrs\brdrw10 \clbrdrb\brdrs\brdrw10 \clbrdrr
\brdrs\brdrw10 \cltxlrtb \cellx6237\pard\plain \s15\widctlpar\intbl\adjustright \f2\fs20\lang2057\cgrid {\b\f0\fs24 Source Texts\cell }{\f0\fs24 5\cell 1\cell 6\cell }\pard\plain \widctlpar\intbl\adjustright \lang2057\cgrid {\row }\pard\plain 
\s15\widctlpar\intbl\adjustright \f2\fs20\lang2057\cgrid {\b\f0\fs24 Translations\cell }{\f0\fs24 1\cell 6\cell 7\cell }\pard\plain \widctlpar\intbl\adjustright \lang2057\cgrid {\row }\trowd \trqc\trgaph108\trrh262\trleft-108\trbrdrt\brdrs\brdrw10 
\trbrdrl\brdrs\brdrw10 \trbrdrb\brdrs\brdrw10 \trbrdrr\brdrs\brdrw10 \trbrdrh\brdrs\brdrw10 \trbrdrv\brdrs\brdrw10 \clvertalt\clbrdrt\brdrs\brdrw10 \clbrdrl\brdrs\brdrw10 \clbrdrb\brdrs\brdrw10 \clbrdrr\brdrs\brdrw10 \cltxlrtb \cellx1843\clvertalt\clbrdrt
\brdrs\brdrw10 \clbrdrl\brdrs\brdrw10 \clbrdrb\brdrs\brdrw10 \clbrdrr\brdrs\brdrw10 \cltxlrtb \cellx3402\clvertalt\clbrdrt\brdrs\brdrw10 \clbrdrl\brdrs\brdrw10 \clbrdrb\brdrs\brdrw10 \clbrdrr\brdrs\brdrw10 \cltxlrtb \cellx5103\clvertalt\clbrdrt
\brdrs\brdrw10 \clbrdrl\brdrs\brdrw10 \clbrdrb\brdrs\brdrw10 \clbrdrr\brdrs\brdrw10 \cltxlrtb \cellx6237\pard\plain \s15\widctlpar\intbl\adjustright \f2\fs20\lang2057\cgrid {\b\f0\fs24 Words\cell }{\f0\fs24 62,039\cell 72,402\cell 134,441\cell 
}\pard\plain \widctlpar\intbl\adjustright \lang2057\cgrid {\row }\pard\plain \s15\widctlpar\adjustright \f2\fs20\lang2057\cgrid {\f0\fs24  
\par 
\par }\pard \s15\widctlpar\outlinelevel0\adjustright {\b\f0\fs24 Encoding Aims and Options
\par }\pard \s15\widctlpar\adjustright {\f0\fs24 The overriding aim of all text encodi
ng options adopted in COMPARA was to provide accurate examples of how sentences have been translated from Portuguese into English and from English into Portuguese, and, within those sentences, to provide "co-textualized" examples of how words and phrases 
have been translated. 
\par 
\par Because COMPARA is a sentence-driven corpus, text divisions that lie above the level of the sentence - such as chapter and paragraph divisions - were not encoded. In fact, since COMPARA did not acquire permission for the actual physi
cal redistribution of texts, no attempt was made to preserve the texts in a format that would allow exact future replication. Page layout, typeface, pictures, diagrams and all other material that is not immediately relevant to the study of language contra
st and translation were simply removed without being replaced by omit tags.  
\par 
\par Also, since the texts in COMPARA can only be used in COMPARA, one did not feel obliged to follow the Text Encoding Initiative (TEI) guidelines (Sperberg-McQueen & Burnard, 1994) or any
 other standard for text encoding to the letter, although an attempt was made to follow the TEI's spirit and syntax whenever dealing with phenomena coped with by the TEI.
\par 
\par }\pard \s15\widctlpar\outlinelevel0\adjustright {\b\f0\fs24 Criteria for Text Alignment
\par }\pard \s15\widctlpar\adjustright {\f0\fs24 The basic unit of alignment in COMPARA is the source-tex
t sentence. Whenever there is not a one-to-one sentence correspondence between source and translation, it is the translation that has been split or joined up to conform to the way sentences were originally divided in the source text. 
\par 
\par Thus an alignment un
it is always one orthographic sentence in the source text and the corresponding text in the translation(s), whether it is one, more than one or even only part of a sentence. Source-text sentences that have been left out of the translation are aligned with
 
blank units. Sentences that have been added to the translation with no corresponding text in the original are fitted into the nearest preceding alignment unit. Sentences that have been reordered in the translation are aligned with the sentences that promp
ted them in the source texts. Thus if the order of sentences A, B and C in the source text has been changed to A, C and B in the translation, A is aligned with A, B is brought back so that it can be aligned with B, and C is aligned with C. Table 2 sum
marizes these alignment criteria.
\par 
\par }\pard \s15\widctlpar\outlinelevel0\adjustright {\f0\fs24\ul \page Table 2: Alignment criteria
\par 
\par }\trowd \trgaph108\trleft-108\trbrdrt\brdrs\brdrw10 \trbrdrl\brdrs\brdrw10 \trbrdrb\brdrs\brdrw10 \trbrdrr\brdrs\brdrw10 \trbrdrh\brdrs\brdrw10 \trbrdrv\brdrs\brdrw10 \clvertalt\clbrdrt\brdrs\brdrw10 \clbrdrl\brdrs\brdrw10 \clbrdrb\brdrs\brdrw10 \clbrdrr
\brdrs\brdrw10 \cltxlrtb \cellx2844\clvertalt\clbrdrt\brdrs\brdrw10 \clbrdrl\brdrs\brdrw10 \clbrdrb\brdrs\brdrw10 \clbrdrr\brdrs\brdrw10 \cltxlrtb \cellx5796\clvertalt\clbrdrt\brdrs\brdrw10 \clbrdrl\brdrs\brdrw10 \clbrdrb\brdrs\brdrw10 \clbrdrr
\brdrs\brdrw10 \cltxlrtb \cellx8748\pard\plain \widctlpar\intbl\adjustright \lang2057\cgrid {\cell SOURCE TEXT \cell TRANSLATION\cell }\pard \widctlpar\intbl\adjustright {\row }\trowd \trgaph108\trleft-108\trbrdrt\brdrs\brdrw10 \trbrdrl\brdrs\brdrw10 
\trbrdrb\brdrs\brdrw10 \trbrdrr\brdrs\brdrw10 \trbrdrh\brdrs\brdrw10 \trbrdrv\brdrs\brdrw10 \clvertalt\clbrdrt\brdrs\brdrw10 \clbrdrl\brdrs\brdrw10 \clbrdrb\brdrs\brdrw10 \clbrdrr\brdrs\brdrw10 \cltxlrtb \cellx2844\clvertalt\clbrdrt\brdrs\brdrw10 \clbrdrl
\brdrs\brdrw10 \clbrdrb\brdrs\brdrw10 \clbrdrr\brdrs\brdrw10 \cltxlrtb \cellx5796\clvertalt\clbrdrt\brdrs\brdrw10 \clbrdrl\brdrs\brdrw10 \clbrdrb\brdrs\brdrw10 \clbrdrr\brdrs\brdrw10 \cltxlrtb \cellx8748\pard\plain 
\s6\keepn\widctlpar\intbl\outlinelevel5\adjustright \i\lang2057\cgrid {Sentence preserved\cell }\pard\plain \qc\widctlpar\intbl\adjustright \lang2057\cgrid {1\cell 1\cell }\pard \widctlpar\intbl\adjustright {\row }\pard\plain 
\s6\keepn\widctlpar\intbl\outlinelevel5\adjustright \i\lang2057\cgrid {Sentence split\cell }\pard\plain \qc\widctlpar\intbl\adjustright \lang2057\cgrid {1\cell 2 \cell }\pard \widctlpar\intbl\adjustright {\row }\pard\plain 
\s6\keepn\widctlpar\intbl\outlinelevel5\adjustright \i\lang2057\cgrid {Sentence joined\cell }\pard\plain \qc\widctlpar\intbl\adjustright \lang2057\cgrid {1\cell }{{\field{\*\fldinst SYMBOL 49 \\f "MS Reference 2" \\s 12}{\fldrslt\f45\fs24}}}{\cell }\pard 
\widctlpar\intbl\adjustright {\row }\pard\plain \s6\keepn\widctlpar\intbl\outlinelevel5\adjustright \i\lang2057\cgrid {Sentence deleted\cell }\pard\plain \qc\widctlpar\intbl\adjustright \lang2057\cgrid {1\cell 0\cell }\pard \widctlpar\intbl\adjustright {
\row }\pard\plain \s6\keepn\widctlpar\intbl\outlinelevel5\adjustright \i\lang2057\cgrid {Sentence added\cell }\pard\plain \qc\widctlpar\intbl\adjustright \lang2057\cgrid {1\cell 1 + [1]\cell }\pard \widctlpar\intbl\adjustright {\row }\trowd 
\trgaph108\trrh691\trleft-108\trbrdrt\brdrs\brdrw10 \trbrdrl\brdrs\brdrw10 \trbrdrb\brdrs\brdrw10 \trbrdrr\brdrs\brdrw10 \trbrdrh\brdrs\brdrw10 \trbrdrv\brdrs\brdrw10 \clvertalt\clbrdrt\brdrs\brdrw10 \clbrdrl\brdrs\brdrw10 \clbrdrb\brdrs\brdrw10 \clbrdrr
\brdrs\brdrw10 \cltxlrtb \cellx2844\clvertalt\clbrdrt\brdrs\brdrw10 \clbrdrl\brdrs\brdrw10 \clbrdrb\brdrs\brdrw10 \clbrdrr\brdrs\brdrw10 \cltxlrtb \cellx5796\clvertalt\clbrdrt\brdrs\brdrw10 \clbrdrl\brdrs\brdrw10 \clbrdrb\brdrs\brdrw10 \clbrdrr
\brdrs\brdrw10 \cltxlrtb \cellx8748\pard\plain \s6\keepn\widctlpar\intbl\outlinelevel5\adjustright \i\lang2057\cgrid {\lang1024 
{\shp{\*\shpinst\shpleft6992\shptop298\shpright7704\shpbottom586\shpfhdr0\shpbxcolumn\shpbypara\shpwr3\shpwrk0\shpfblwtxt0\shpz0\shplid1026{\sp{\sn shapeType}{\sv 105}}{\sp{\sn fFlipH}{\sv 0}}{\sp{\sn fFlipV}{\sv 0}}
{\sp{\sn rotation}{\sv -11554946}}{\sp{\sn adjustValue}{\sv 11182}}{\sp{\sn adjust2Value}{\sv 17408}}}{\shprslt{\*\do\dobxcolumn\dobypara\dodhgt8192\dppolygon\dppolycount85\dpptx0\dppty288\dpptx1\dppty259\dpptx6\dppty230\dpptx11\dppty202\dpptx21\dppty176
\dpptx32\dppty151\dpptx44\dppty127\dpptx59\dppty105\dpptx67\dppty94\dpptx76\dppty84\dpptx85\dppty75\dpptx95\dppty66\dpptx104\dppty57\dpptx114\dppty49\dpptx124\dppty42\dpptx136\dppty35\dpptx147\dppty29\dpptx158\dppty23\dpptx170\dppty18\dpptx182\dppty13
\dpptx194\dppty9\dpptx206\dppty6\dpptx220\dppty3\dpptx232\dppty1\dpptx245\dppty0\dpptx259\dppty0\dpptx327\dppty0\dpptx347\dppty1\dpptx367\dppty4\dpptx387\dppty8\dpptx406\dppty14\dpptx425\dppty21\dpptx442\dppty30\dpptx460\dppty41\dpptx476\dppty53
\dpptx492\dppty66\dpptx507\dppty81\dpptx521\dppty97\dpptx533\dppty113\dpptx545\dppty132\dpptx555\dppty151\dpptx563\dppty171\dpptx571\dppty192\dpptx712\dppty192\dpptx551\dppty288\dpptx361\dppty192\dpptx502\dppty192\dpptx496\dppty173\dpptx489\dppty155
\dpptx480\dppty138\dpptx469\dppty122\dpptx459\dppty106\dpptx448\dppty91\dpptx435\dppty77\dpptx421\dppty65\dpptx408\dppty53\dpptx393\dppty42\dpptx378\dppty32\dpptx362\dppty24\dpptx345\dppty17\dpptx328\dppty11\dpptx311\dppty6\dpptx293\dppty2\dpptx293\dppty2
\dpptx281\dppty5\dpptx269\dppty7\dpptx257\dppty11\dpptx246\dppty14\dpptx235\dppty18\dpptx224\dppty24\dpptx214\dppty29\dpptx204\dppty35\dpptx194\dppty41\dpptx183\dppty48\dpptx165\dppty63\dpptx148\dppty80\dpptx132\dppty98\dpptx118\dppty118\dpptx106\dppty139
\dpptx95\dppty162\dpptx85\dppty185\dpptx79\dppty210\dpptx73\dppty235\dpptx69\dppty261\dpptx68\dppty288\dpx6992\dpy298\dpxsize712\dpysize288
\dpfillfgcr255\dpfillfgcg255\dpfillfgcb255\dpfillbgcr255\dpfillbgcg255\dpfillbgcb255\dpfillpat1\dplinew15\dplinecor0\dplinecog0\dplinecob0}}}}{Sentence reordered\cell }\pard\plain \qc\widctlpar\intbl\adjustright \lang2057\cgrid {A, B, C\cell A,    C,    B
\cell }\pard \widctlpar\intbl\adjustright {\row }\pard\plain \s15\widctlpar\outlinelevel0\adjustright \f2\fs20\lang2057\cgrid {\f0\fs24\ul 
\par }\pard \s15\widctlpar\adjustright {\f0\fs24 
\par For all the above cases, special alignment markup is inserted so that corpus users can search for translational discourse changes such as when and where translators have chosen to join, split, delete, add or reorder sentences in the translation. 
\par 
\par On the one han
d, it is important to note that the alignment markup in COMPARA does not capture the addition or deletion or reordering of units smaller than the sentence such as individual words, phrases and clauses. On the other hand, the strength of this alignment pro
c
edure is that it enables one to align a source text with more than one translation, and to compare not only source and target text, but also more than one translation of the same source (where the source would, in this case, act as a common denominator to
 the two translations).  
\par 
\par 
\par }\pard \s15\widctlpar\outlinelevel0\adjustright {\b\f0\fs24 Text Preparation: from print to Web
\par }\pard \s15\widctlpar\adjustright {\f0\fs24 
\par The procedure for preparing texts for COMPARA is as follows:
\par 
\par {\pntext\pard\plain\s15 \lang2057\cgrid \hich\af0\dbch\af0\loch\f0 1.\tab}}\pard \s15\fi-360\li360\widctlpar\jclisttab\tx360{\*\pn \pnlvlbody\ilvl0\ls7\pnrnot0\pndec\pnstart1\pnindent360\pnhang{\pntxta .}}\ls7\adjustright {\f0\fs24 
The texts in the corpus that are not available in electronic form are scanned and submitted to an optical character recognition (OCR) program.
\par {\pntext\pard\plain\s15 \lang2057\cgrid \hich\af0\dbch\af0\loch\f0 2.\tab}}\pard \s15\fi-360\li360\widctlpar\jclisttab\tx360{\*\pn \pnlvlbody\ilvl0\ls7\pnrnot0\pndec\pnstart1\pnindent360\pnhang{\pntxta .}}\ls7\adjustright {\f0\fs24 
The OCR is revised (if the text was scanned), all non-translational material is removed, and marks for titles, foreign words and expressions, emphasis and translators' notes are introduced. The notes themselves are inserted at the point
 where their identifiers appear.
\par {\pntext\pard\plain\s15 \lang2057\cgrid \hich\af0\dbch\af0\loch\f0 3.\tab}}\pard \s15\fi-360\li360\widctlpar\jclisttab\tx360{\*\pn \pnlvlbody\ilvl0\ls7\pnrnot0\pndec\pnstart1\pnindent360\pnhang{\pntxta .}}\ls7\adjustright {\f0\fs24 
Source text and translation are aligned manually, paragraph by paragraph.
\par {\pntext\pard\plain\s15 \lang2057\cgrid \hich\af0\dbch\af0\loch\f0 4.\tab}}\pard \s15\fi-360\li360\widctlpar\jclisttab\tx360{\*\pn \pnlvlbody\ilvl0\ls7\pnrnot0\pndec\pnstart1\pnindent360\pnhang{\pntxta .}}\ls7\adjustright {\f0\fs24 
The texts are submitted to a program of automatic tokenization and sentence separation developed by the AC/DC project  (Santos et al. 2000) and to a program of automatic sentence alignment - the  IMS Corpus Workbench Easyalign. 
\par {\pntext\pard\plain\s15 \lang2057\cgrid \hich\af0\dbch\af0\loch\f0 5.\tab}}\pard \s15\fi-360\li360\widctlpar\jclisttab\tx360{\*\pn \pnlvlbody\ilvl0\ls7\pnrnot0\pndec\pnstart1\pnindent360\pnhang{\pntxta .}}\ls7\adjustright {\f0\fs24 
The alignment results are revised semi-automatically so as to conform to COMPARA's alignment criteria. Alignment markup for sentence joining and reordering is introduced manually at this point. 
\par {\pntext\pard\plain\s15 \lang2057\cgrid \hich\af0\dbch\af0\loch\f0 6.\tab}}\pard \s15\fi-360\li360\widctlpar\jclisttab\tx360{\*\pn \pnlvlbody\ilvl0\ls7\pnrnot0\pndec\pnstart1\pnindent360\pnhang{\pntxta .}}\ls7\adjustright {\f0\fs24 
The remaining alignment markup (for sentence deletion, addition and splitting) is inserted automatically.
\par {\pntext\pard\plain\s15 \lang2057\cgrid \hich\af0\dbch\af0\loch\f0 7.\tab}}\pard \s15\fi-360\li360\widctlpar\jclisttab\tx360{\*\pn \pnlvlbody\ilvl0\ls7\pnrnot0\pndec\pnstart1\pnindent360\pnhang{\pntxta .}}\ls7\adjustright {\f0\fs24 
Alignment markup revision (sentence addition markup has to be discriminated from sentence splitting markup manually). 
\par {\pntext\pard\plain\s15 \lang2057\cgrid \hich\af0\dbch\af0\loch\f0 8.\tab}}\pard \s15\fi-360\li360\widctlpar\jclisttab\tx360{\*\pn \pnlvlbody\ilvl0\ls7\pnrnot0\pndec\pnstart1\pnindent360\pnhang{\pntxta .}}\ls7\adjustright {\f0\fs24 
IMS-Corpus Workbench automatic encoding}{\f0\fs24\super 5}{\f0\fs24 .
\par }\pard \s15\widctlpar\adjustright {\f0\fs24 
\par 
\par 
\par }\pard \s15\widctlpar\outlinelevel0\adjustright {\b\f0\fs24 
\par Searching COMPARA: the DISPARA interface 
\par }\pard \s15\widctlpar\adjustright {\f0\fs24 COMPARA can be searched via the DISPARA interface, which has been developed as a bridge between the IMS Corpus Workbench system used and the specific requirements 
of COMPARA. Although DISPARA was conceived to cater for the specific needs of COMPARA, it can also be very easily adapted to other parallel corpora that are encoded according to the IMS Corpus Workbench system}{\f0\fs24\super 6}{\f0\fs24 . 
\par  
\par Two search options are available in DISPA
RA. The Simple Search was made for people who have never used a corpus before. It allows users to search the entire corpus either in the Portuguese-English or in the English-Portuguese direction. The instructions on how to conduct a Simple Search are extr
emely simple. Users only have to write a word or expression in English or Portuguese and press the search button. No special training is required (see appendix 1).
\par 
\par The Complex Search was made for those who find the Simple Search too restrictive and want to
 conduct more sophisticated queries. We have endeavoured to make the Complex Search as user-friendly as possible, so that people who have never used a corpus before should feel confident enough to exploit its potentialities. A user-testing session is to t
ake place shortly so as to provide us with feedback on what can be improved in the Complex Search option. As it stands, users are guided through four relatively simple search steps (see appendix 2).
\par 
\par }\pard \s15\widctlpar\outlinelevel0\adjustright {\f0\fs24\ul Step 1
\par }\pard \s15\widctlpar\adjustright {\f0\fs24 In step one, users are asked to choose their searc
h direction. As in the Simple Search, they can search from Portuguese to English or from English to Portuguese. However, in the Complex Search, instead of searching the whole corpus, users can also tell the system that they only want to search from source
-texts to translations, or only from translations to source texts. The latter is an important option to consider if the directionality of translation is relevant to a particular query. 
\par 
\par }\pard \s15\widctlpar\outlinelevel0\adjustright {\f0\fs24\ul Step 2
\par }\pard \s15\widctlpar\adjustright {\f0\fs24 In step two, users are asked if they want to narrow down the co
rpus, and, if so, they are asked to choose which texts within the corpus you want to use. This is a very important step because, as COMPARA is an open-ended corpus, it is here that users will be able to control which texts they are going to use if their q
ueries require a balanced corpus or a specific subset or other of the corpus. 
\par 
\par COMPARA can be automatically narrowed down so as to search only within specific varieties of Portuguese and English. It is possible to select any combination of Portuguese and E
nglish language varieties. For example, users can tell DISPARA that they want to search only Brazilian Portuguese and British English, or all varieties of Portuguese but only American English, etc. 
\par 
\par Next, it is possible to narrow down the corpus by date of publication. Users who are not interested in non-contemporary language, for example, can automatically remove source texts and translations published before a particular date. 
\par 
\par The third narrowing-down option available allows users to select any manual 
combination of texts. Users can tell DISPARA exactly which texts they want to use for their search queries, and create their own, tailor-made sub-corpus of COMPARA. They are thus able to conduct searches within texts by only one particular author, or grou
p of authors, or translator, and so on. 
\par 
\par Eventually, when other genres are added to the corpus, there will also be an option that allows users to select texts automatically by genre. 
\par 
\par }\pard \s15\widctlpar\outlinelevel0\adjustright {\f0\fs24\ul Step 3
\par }\pard \s15\widctlpar\adjustright {\f0\fs24 The third step of the Complex Search enables users to select dif
ferent displays of the results. Users can inspect concordances, distribution of forms, distribution of sources (how a search expression is distributed in the texts within the corpus) and a quantitative wrap up (the distribution of the search expression in
 the two languages, for searches that involve alignment constraints - see below).
\par 
\par }\pard \s15\widctlpar\outlinelevel0\adjustright {\f0\fs24\ul Step 4
\par }\pard \s15\widctlpar\adjustright {\f0\fs24 In the fourth and final step of the Complex Search, users are asked to enter their search queries. The IMS Corpus Workbench syntax can be used here to refine searches 
so as to include in a single query access to different spellings of a word (for example, analyse and analyze), different morphological variants of a word (for example, walk, walked, walks, etc.), a word and a collocate with any number of elements in betwe
en, and so on}{\f0\fs24\super 7}{\f0\fs24 
. Although for now users still need to learn to use the IMS Corpus Workbench syntax to have access to those details, DISPARA has plans to develop a more user-friendly interface for the types of queries that prove to be more popular among users.
\par 
\par Apart from entering a given search word or expression, in the Complex Search users can also enter an alignment constraint. For example, users searching for the Portuguese translation of }{\i\f0\fs24 yes}{\f0\fs24 , which is usually rendered as }{
\i\f0\fs24 sim}{\f0\fs24 , can retrieve just the cases in which }{\i\f0\fs24 yes}{\f0\fs24  is translated into }{\i\f0\fs24 sim}{\f0\fs24  or just the cases in which }{\i\f0\fs24 yes}{\f0\fs24  is translated into something other than }{\i\f0\fs24 sim}{
\f0\fs24 . 
\par 
\par Some searchable features that are very specific to COMPARA are already directly available through the DISPARA interface. Whatever the query, D
ISPARA allows users to inspect translators' notes and alignment properties. In adddition to this, users can search directly for translators' notes, emphasis, foreign words and expressions, and titles. And because of the way the texts in COMPARA have been 
a
ligned and encoded, it is also possible to inspect when and where translators have decided to join, separate, delete and add sentences to the translation. The possibility of inspecting reordered sentences was not yet operational at the time this paper was
 written. 
\par 
\par }\pard \s15\widctlpar\outlinelevel0\adjustright {\b\f0\fs24 Search results
\par }\pard \s15\widctlpar\adjustright {\f0\fs24 The users of COMPARA are welcome to use the results of their search queries for research and education. 
\par 
\par The maximum number of concordances per query shown is 500, because many of the texts in COMPARA are still in copyright. If 
the corpus user chooses to narrow down the corpus in any way (e.g. by language variety, by date of publication, or by selecting a specific text pair), then the maximum number of concordances per query is further limited to 200. Whenever the results exceed
 these numbers, a random selection of respectively 500 and 200 concordances are presented instead. However, even when it is not possible to }{\i\f0\fs24 show}{\f0\fs24 
 all the concordances, the total number of solutions found is always given. Thus the user will get a message saying that, for copyright reasons, only 500 (or 200) random concordances out of the x>500 (or x>200) found can be presented. 
\par 
\par The concordances are displayed in two vertical columns, with the Portuguese or English search item appearing in bold on the left-hand side, and the corresponding text in English or Portuguese on the right-hand side. Instead of a key-word-in-context (KWIC
) concordance with a fixed number of characters to the left and to the right, the user gets a KWIC concordance where the context is one full source-text sen
tence and the corresponding text in the translation (see appendix 3). There are plans to allow the user to expand the amount of co-text given within the limits of fair-use, but this feature was not yet operational at the time this paper was written.  

\par 
\par Nex
t to each parallel concordance displayed, there is a link to the full reference of the pair of texts from where the parallel concordance was retrieved. When looking up a reference, users also get information on copyright, on language variety, and on the n
umber of words and alignment units for the extracts in question. 
\par 
\par It is possible to scroll up and down the results screen to see all the concordances displayed, and it is possible to save the results in html, text or even to cut and paste them into a word-processing program. 
\par 
\par }\pard \s15\widctlpar\outlinelevel0\adjustright {\b\f0\fs24 Conclusion
\par }\pard \s15\widctlpar\adjustright {\f0\fs24 
The COMPARA project began only a year ago and is still very much in its infancy. There is a long list of fiction texts for which copyright clearance has been acquired, but which are still waiting to be processed. There are 
plans to expand the corpus so as to include genres other than fiction, but this phase has not yet begun. The DISPARA interface to COMPARA has reached a reasonable stage of development, but feedback from users is essential to further its objectives of beco
m
ing truly user-friendly and accessible to all. This is just the beginning, and we hope some of the innovative choices made during the process of creating COMPARA and DISPARA will contribute towards future developments in the conception and distribution of
 parallel corpora.
\par }\pard \s15\widctlpar\outlinelevel0\adjustright {\b\f0\fs24 
\par Notes
\par }\pard \s15\widctlpar\adjustright {\f0 
1  Parallel is being used here to refer to a bilingual collection of source texts and their translations. In the contrastive linguistics tradition, this would have been referred to as a translation corpus. Johansson (1998) predicte
d that the problem of conflicting terminology would eventually be resolved as the field developed and usage became more settled. This does not seem to have happened yet. See also the introduction of V\'e9ronis (2000).
\par 2 The Computational Processing of Portugu
ese project is concerned with the creation, evaluation, cataloguing and public distribution of Portuguese language computational resources. It is financed by the Portuguese Ministry of Science and Technology. For further information, see Santos (2000) and
 http://www.portugues.mct.pt. 
\par 3 Unlike the texts in the English-Norwegian Parallel Corpus, the extracts in COMPARA are not required to be of about the same length nor taken only from the beginning of novels. The ENPC\rquote s attempt to achieve homogeneity in thi
s respect was found to be a disadvantage by Santos and Oskefjell (1999) when attempting to validate corpus-based contrastive work.
\par 4 For a full regularly updated list of copyright permissions, see http://www.portugues.mct.pt/COMPARA/CorpusContents.html.
\par 5 
The IMS Corpus Workbench (Christ, 1994; Christ et al. 1999) has been singled out in previous occasions as the best corpus system available given the general context of the Computational Processing of Portuguese project, which is responsible for creating t
he DISPARA Web interface to COMPARA. Motivations for its use can be found in Santos (1998), and Santos & Ranchhod (1999).
\par 6 DISPARA is a general system for DIStributing PARAlell corpora on the Web. 
\par 7 For a detailed description of the options available, see the IMS-CPQ User's Manual at http://www.ims.uni-stutgart.de/CorpusWorkbench/
\par 
\par }\pard \s15\widctlpar\outlinelevel0\adjustright {\b\f0\fs24 References
\par }\pard \s15\widctlpar\adjustright {\f0 Alencar, Jos\'e9 de (1865) }{\i\f0 Iracema,}{\f0  http://www.vbookstore.com.br/nacional/josedealencar/iracema.shtml [06/12/1999]. Digital text prepared by the Biblioteca Virtual do E
studante Brasileiro, based on 24th ed.,S\'e3o Paulo: \'c1tica, 1991
\par 
\par ----- (1886) }{\i\f0 Iracema, the honey lips: a legend of Brazil,}{\f0  London: Bickers. 
\par 
\par ----- (2000) }{\i\f0 Iracema}{\f0 , New York: Oxford University Press.
\par  
\par Christ, Oliver, B. Schulze, A. Hofmann &  E. Koenig (1999) "The IMS Corpus Workbench: Corpus Query Processor (CQP): User's Manual", Institute for Natural Language Processing, University of Stuttgart, March 8, 1999 (CQP V2.2). 
\par 
\par Johansson, Stig (1998) "On the role of  corpora in cross-linguistic research" in S. Johansson & S. Oksefjell (eds)}{\i\f0  Corpora and crosslinguistic research: theory, method and case studies}{\f0 , Amsterdam: Rodopi, pp 3-24.
\par Johansson, Stig, J. Ebeling & S. Oksefjell (1999) English-Norwegian Parallel Corpus: Manual http://www.hf.uio.no/iba/prosjekt/ENPCmanual.html [Access Date 7/7/2000]
\par }\pard \s15\widctlpar\outlinelevel0\adjustright {\f0 
\par }\pard \s15\widctlpar\adjustright {\f0 Lodge, David (1975) }{\i\f0 Changing Places}{\f0 , London: Secker & Warburg.
\par 
\par ----- (1995a) }{\i\f0 Therapy}{\f0 , London: Secker & Warburg.
\par 
\par ----- (1995b) }{\i\f0 A Troca}{\f0  [Changing Places], Porto: Asa.
\par 
\par ----- (1995c) }{\i\f0 Terapia}{\f0  [Therapy], Lisboa: Gradiva. 
\par 
\par ----- (1997) }{\i\f0 Terapia}{\f0  [Therapy], S\'e3o Paulo: Scipione. 
\par 
\par ----- (1998) }{\i\f0 Invertendo os Pap\'e9is}{\f0  [Changing Places], S\'e3o Paulo: Scipione.
\par 
\par Santos, Diana (1998) "Providing access to language resources through the World Wide Web: the Oslo Corpus of Bosnian Texts", in A.Rubio, N.Gallardo, R.Castro and A.Tejada (eds) }{\i\f0 
Proceedings of The First International Conference on Language Resources and Evaluation, }{\f0 Vol. 1, pp.475-481.
\par 
\par ------ & S. Oksefjell (1999) "Using a Parallel Corpus to Validate
\par Independent Claims", }{\i\f0 Languages in Contrast}{\f0 , Vol. 2(1):117-132. 
\par 
\par ----- & E. Ranchhod (1999) "Ambientes de processamento de corpora em portugu\'eas: Compara\'e7\'e3o entre dois sistemas", in }{\i\f0 Actas do IV Encontro sobre o Processamento Computacional da L\'edngua Portuguesa}{\f0  (Escrita e Falada), PROPOR [
\ldblquote Portuguese corpora processing: comparing two systems\rdblquote , in }{\i\f0 Proceedings of the IV Encounter on the computational processing of written and spoken Portuguese}{\f0 ] (\'c9vora, 20-21 September 1999), pp. 257-268.
\par 
\par ----- & E. Bick (2000) "Providing Internet access to Portuguese corpora: the AC/DC project", in Gavriladou M., G. Carayannis, S. Markantonatou, S.Piperidis & G. Stainhaouer (eds) }{\i\f0 
Proceedings of the Second International Conference on Language Resources and Evaluation, LREC2000}{\f0 , pp.205-210. 
\par 
\par ----- (2000) "O projecto Processamento Computacional do Portugu\'eas: Balan\'e7o e perspectivas", in M. Gra\'e7a Nunes (ed) }{\i\f0 Actas do V Encontro para o processamento computacional da l\'edngua portuguesa escrita e falada}{\f0 
 (PROPOR 2000) [The Computational Processing of Portuguese project: balance and perspectives, in }{\i\f0 Proceedings of the V Encounter on the computational processing of written and spoken Portuguese}{\f0 ], pp.105-113.
\par 
\par Sperberg-McQueen, C. & Burnard, L. (eds) (1994) \ldblquote Guidelines for Electronic Text Encoding and Interchange\rdblquote  }{\i\f0 TEI P3}{\f0 
. Association for Computers and Humanities/ Association for Computational Linguistics/ Association for Literary and Linguistic Computing. Chicago and Oxford. 
\par }\pard \s15\widctlpar\outlinelevel0\adjustright {\f0 
\par }\pard\plain \widctlpar\adjustright \lang2057\cgrid {\fs20 V\'e9ronis, Jean (ed) (2000) }{\i\fs20 Parallel Text Processing}{\fs20 , Dordrecht: Kluwer Academic Publishers.
\par }}