OXIESEC PANEL
- Current Dir:
/
/
opt
/
alt
/
libicu71
/
usr
/
include
/
unicode
Server IP: 2a02:4780:11:1594:0:ef5:22d7:a
Upload:
Create Dir:
Name
Size
Modified
Perms
📁
..
-
05/14/2024 03:54:45 PM
rwxr-xr-x
📄
alphaindex.h
26.52 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
appendable.h
8.53 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
basictz.h
9.86 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
brkiter.h
27.81 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
bytestream.h
10.75 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
bytestrie.h
20.77 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
bytestriebuilder.h
7.46 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
calendar.h
105.88 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
caniter.h
7.44 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
casemap.h
25.33 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
char16ptr.h
7.22 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
chariter.h
24.06 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
choicfmt.h
23.97 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
coleitr.h
13.77 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
coll.h
56.26 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
compactdecimalformat.h
6.88 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
curramt.h
3.78 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
currpinf.h
7.3 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
currunit.h
4.02 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
datefmt.h
40.71 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
dbbi.h
1.19 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
dcfmtsym.h
20.61 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
decimfmt.h
87.57 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
docmain.h
7.21 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
dtfmtsym.h
38.21 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
dtintrv.h
3.84 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
dtitvfmt.h
49.27 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
dtitvinf.h
18.63 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
dtptngen.h
28.84 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
dtrule.h
8.69 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
edits.h
20.74 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
enumset.h
2.08 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
errorcode.h
4.84 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
fieldpos.h
8.7 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
filteredbrk.h
5.37 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
fmtable.h
24.43 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
format.h
12.5 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
formattedvalue.h
9.75 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
fpositer.h
3.03 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
gender.h
3.33 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
gregocal.h
31.88 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
icudataver.h
1.02 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
icuplug.h
11.83 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
idna.h
12.7 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
listformatter.h
8.79 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
localebuilder.h
11.08 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
localematcher.h
26.84 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
localpointer.h
19.36 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
locdspnm.h
7.12 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
locid.h
47.66 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
measfmt.h
11.42 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
measunit.h
104.97 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
measure.h
4.33 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
messagepattern.h
33.71 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
msgfmt.h
44.22 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
normalizer2.h
33.66 KB
12/01/2022 11:12:32 AM
rw-r--r--
📄
normlzr.h
30.95 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
nounit.h
2.25 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
numberformatter.h
94.15 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
numberrangeformatter.h
25.19 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
numfmt.h
49.84 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
numsys.h
7.33 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
parseerr.h
3.08 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
parsepos.h
5.57 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
platform.h
28.09 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
plurfmt.h
25.24 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
plurrule.h
20.82 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ptypes.h
3.49 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
putil.h
6.32 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
rbbi.h
28.64 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
rbnf.h
48.85 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
rbtz.h
15.77 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
regex.h
84.36 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
region.h
9.18 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
reldatefmt.h
22.22 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
rep.h
9.37 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
resbund.h
18.08 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
schriter.h
6.36 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
scientificnumberformatter.h
6.43 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
search.h
22.22 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
selfmt.h
14.34 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
simpleformatter.h
12.59 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
simpletz.h
45.64 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
smpdtfmt.h
71.07 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
sortkey.h
11.18 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
std_string.h
1.05 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
strenum.h
9.92 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
stringoptions.h
5.79 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
stringpiece.h
10.05 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
stringtriebuilder.h
15.47 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
stsearch.h
21.41 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
symtable.h
4.27 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
tblcoll.h
36.92 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
timezone.h
43.81 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
tmunit.h
3.4 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
tmutamt.h
4.91 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
tmutfmt.h
7.42 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
translit.h
65.82 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
tzfmt.h
42.93 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
tznames.h
16.85 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
tzrule.h
35.6 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
tztrans.h
6.13 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ubidi.h
89.61 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ubiditransform.h
12.7 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ubrk.h
24.43 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ucal.h
60.68 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ucasemap.h
15.21 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ucat.h
5.35 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
uchar.h
144.64 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ucharstrie.h
22.53 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ucharstriebuilder.h
7.47 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
uchriter.h
13.42 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
uclean.h
11.2 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ucnv.h
83.46 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ucnv_cb.h
6.58 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ucnv_err.h
20.98 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ucnvsel.h
6.19 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ucol.h
61.95 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ucoleitr.h
9.82 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
uconfig.h
12.07 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ucpmap.h
5.53 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ucptrie.h
22.5 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ucsdet.h
14.69 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ucurr.h
16.72 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
udat.h
62.4 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
udata.h
15.63 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
udateintervalformat.h
11.93 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
udatpg.h
30.18 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
udisplaycontext.h
5.94 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
uenum.h
7.79 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ufieldpositer.h
4.41 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
uformattable.h
10.97 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
uformattedvalue.h
12.3 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ugender.h
2.06 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
uidna.h
33.43 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
uiter.h
22.75 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
uldnames.h
10.48 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ulistformatter.h
10.78 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
uloc.h
54.63 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ulocdata.h
11.3 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
umachine.h
16.05 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
umisc.h
1.33 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
umsg.h
24.25 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
umutablecptrie.h
8.29 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
unifilt.h
4 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
unifunct.h
4.05 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
unimatch.h
6.1 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
unirepl.h
3.38 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
uniset.h
66.46 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
unistr.h
170.53 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
unorm.h
20.55 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
unorm2.h
24.68 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
unounclass.h
735 bytes
12/01/2022 11:12:33 AM
rw-r--r--
📄
unum.h
55.31 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
unumberformatter.h
30.28 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
unumberrangeformatter.h
15.37 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
unumsys.h
7.26 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
uobject.h
10.68 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
upluralrules.h
8.79 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
uregex.h
71.99 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
uregion.h
9.81 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ureldatefmt.h
17.04 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
urename.h
135 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
urep.h
5.38 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ures.h
36.54 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
uscript.h
27.63 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
usearch.h
39.21 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
uset.h
44.29 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
usetiter.h
9.68 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ushape.h
18 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
uspoof.h
65.84 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
usprep.h
8.19 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ustdio.h
38.56 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ustream.h
1.89 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ustring.h
72.36 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
ustringtrie.h
3.15 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
utext.h
58.08 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
utf.h
7.87 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
utf16.h
23.35 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
utf32.h
763 bytes
12/01/2022 11:12:33 AM
rw-r--r--
📄
utf8.h
30.83 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
utf_old.h
45.83 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
utmscale.h
13.78 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
utrace.h
17.18 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
utrans.h
25.54 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
utypes.h
31.06 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
uvernum.h
6.68 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
uversion.h
5.99 KB
12/01/2022 11:12:33 AM
rw-r--r--
📄
vtzone.h
20.74 KB
12/01/2022 11:12:33 AM
rw-r--r--
Editing: ushape.h
Close
// © 2016 and later: Unicode, Inc. and others. // License & terms of use: http://www.unicode.org/copyright.html /* ****************************************************************************** * * Copyright (C) 2000-2012, International Business Machines * Corporation and others. All Rights Reserved. * ****************************************************************************** * file name: ushape.h * encoding: UTF-8 * tab size: 8 (not used) * indentation:4 * * created on: 2000jun29 * created by: Markus W. Scherer */ #ifndef __USHAPE_H__ #define __USHAPE_H__ #include "unicode/utypes.h" /** * \file * \brief C API: Arabic shaping * */ /** * Shape Arabic text on a character basis. * * <p>This function performs basic operations for "shaping" Arabic text. It is most * useful for use with legacy data formats and legacy display technology * (simple terminals). All operations are performed on Unicode characters.</p> * * <p>Text-based shaping means that some character code points in the text are * replaced by others depending on the context. It transforms one kind of text * into another. In comparison, modern displays for Arabic text select * appropriate, context-dependent font glyphs for each text element, which means * that they transform text into a glyph vector.</p> * * <p>Text transformations are necessary when modern display technology is not * available or when text needs to be transformed to or from legacy formats that * use "shaped" characters. Since the Arabic script is cursive, connecting * adjacent letters to each other, computers select images for each letter based * on the surrounding letters. This usually results in four images per Arabic * letter: initial, middle, final, and isolated forms. In Unicode, on the other * hand, letters are normally stored abstract, and a display system is expected * to select the necessary glyphs. (This makes searching and other text * processing easier because the same letter has only one code.) It is possible * to mimic this with text transformations because there are characters in * Unicode that are rendered as letters with a specific shape * (or cursive connectivity). They were included for interoperability with * legacy systems and codepages, and for unsophisticated display systems.</p> * * <p>A second kind of text transformations is supported for Arabic digits: * For compatibility with legacy codepages that only include European digits, * it is possible to replace one set of digits by another, changing the * character code points. These operations can be performed for either * Arabic-Indic Digits (U+0660...U+0669) or Eastern (Extended) Arabic-Indic * digits (U+06f0...U+06f9).</p> * * <p>Some replacements may result in more or fewer characters (code points). * By default, this means that the destination buffer may receive text with a * length different from the source length. Some legacy systems rely on the * length of the text to be constant. They expect extra spaces to be added * or consumed either next to the affected character or at the end of the * text.</p> * * <p>For details about the available operations, see the description of the * <code>U_SHAPE_...</code> options.</p> * * @param source The input text. * * @param sourceLength The number of UChars in <code>source</code>. * * @param dest The destination buffer that will receive the results of the * requested operations. It may be <code>NULL</code> only if * <code>destSize</code> is 0. The source and destination must not * overlap. * * @param destSize The size (capacity) of the destination buffer in UChars. * If <code>destSize</code> is 0, then no output is produced, * but the necessary buffer size is returned ("preflighting"). * * @param options This is a 32-bit set of flags that specify the operations * that are performed on the input text. If no error occurs, * then the result will always be written to the destination * buffer. * * @param pErrorCode must be a valid pointer to an error code value, * which must not indicate a failure before the function call. * * @return The number of UChars written to the destination buffer. * If an error occurred, then no output was written, or it may be * incomplete. If <code>U_BUFFER_OVERFLOW_ERROR</code> is set, then * the return value indicates the necessary destination buffer size. * @stable ICU 2.0 */ U_CAPI int32_t U_EXPORT2 u_shapeArabic(const UChar *source, int32_t sourceLength, UChar *dest, int32_t destSize, uint32_t options, UErrorCode *pErrorCode); /** * Memory option: allow the result to have a different length than the source. * Affects: LamAlef options * @stable ICU 2.0 */ #define U_SHAPE_LENGTH_GROW_SHRINK 0 /** * Memory option: allow the result to have a different length than the source. * Affects: LamAlef options * This option is an alias to U_SHAPE_LENGTH_GROW_SHRINK * @stable ICU 4.2 */ #define U_SHAPE_LAMALEF_RESIZE 0 /** * Memory option: the result must have the same length as the source. * If more room is necessary, then try to consume spaces next to modified characters. * @stable ICU 2.0 */ #define U_SHAPE_LENGTH_FIXED_SPACES_NEAR 1 /** * Memory option: the result must have the same length as the source. * If more room is necessary, then try to consume spaces next to modified characters. * Affects: LamAlef options * This option is an alias to U_SHAPE_LENGTH_FIXED_SPACES_NEAR * @stable ICU 4.2 */ #define U_SHAPE_LAMALEF_NEAR 1 /** * Memory option: the result must have the same length as the source. * If more room is necessary, then try to consume spaces at the end of the text. * @stable ICU 2.0 */ #define U_SHAPE_LENGTH_FIXED_SPACES_AT_END 2 /** * Memory option: the result must have the same length as the source. * If more room is necessary, then try to consume spaces at the end of the text. * Affects: LamAlef options * This option is an alias to U_SHAPE_LENGTH_FIXED_SPACES_AT_END * @stable ICU 4.2 */ #define U_SHAPE_LAMALEF_END 2 /** * Memory option: the result must have the same length as the source. * If more room is necessary, then try to consume spaces at the beginning of the text. * @stable ICU 2.0 */ #define U_SHAPE_LENGTH_FIXED_SPACES_AT_BEGINNING 3 /** * Memory option: the result must have the same length as the source. * If more room is necessary, then try to consume spaces at the beginning of the text. * Affects: LamAlef options * This option is an alias to U_SHAPE_LENGTH_FIXED_SPACES_AT_BEGINNING * @stable ICU 4.2 */ #define U_SHAPE_LAMALEF_BEGIN 3 /** * Memory option: the result must have the same length as the source. * Shaping Mode: For each LAMALEF character found, expand LAMALEF using space at end. * If there is no space at end, use spaces at beginning of the buffer. If there * is no space at beginning of the buffer, use spaces at the near (i.e. the space * after the LAMALEF character). * If there are no spaces found, an error U_NO_SPACE_AVAILABLE (as defined in utypes.h) * will be set in pErrorCode * * Deshaping Mode: Perform the same function as the flag equals U_SHAPE_LAMALEF_END. * Affects: LamAlef options * @stable ICU 4.2 */ #define U_SHAPE_LAMALEF_AUTO 0x10000 /** Bit mask for memory options. @stable ICU 2.0 */ #define U_SHAPE_LENGTH_MASK 0x10003 /* Changed old value 3 */ /** * Bit mask for LamAlef memory options. * @stable ICU 4.2 */ #define U_SHAPE_LAMALEF_MASK 0x10003 /* updated */ /** Direction indicator: the source is in logical (keyboard) order. @stable ICU 2.0 */ #define U_SHAPE_TEXT_DIRECTION_LOGICAL 0 /** * Direction indicator: * the source is in visual RTL order, * the rightmost displayed character stored first. * This option is an alias to U_SHAPE_TEXT_DIRECTION_LOGICAL * @stable ICU 4.2 */ #define U_SHAPE_TEXT_DIRECTION_VISUAL_RTL 0 /** * Direction indicator: * the source is in visual LTR order, * the leftmost displayed character stored first. * @stable ICU 2.0 */ #define U_SHAPE_TEXT_DIRECTION_VISUAL_LTR 4 /** Bit mask for direction indicators. @stable ICU 2.0 */ #define U_SHAPE_TEXT_DIRECTION_MASK 4 /** Letter shaping option: do not perform letter shaping. @stable ICU 2.0 */ #define U_SHAPE_LETTERS_NOOP 0 /** Letter shaping option: replace abstract letter characters by "shaped" ones. @stable ICU 2.0 */ #define U_SHAPE_LETTERS_SHAPE 8 /** Letter shaping option: replace "shaped" letter characters by abstract ones. @stable ICU 2.0 */ #define U_SHAPE_LETTERS_UNSHAPE 0x10 /** * Letter shaping option: replace abstract letter characters by "shaped" ones. * The only difference with U_SHAPE_LETTERS_SHAPE is that Tashkeel letters * are always "shaped" into the isolated form instead of the medial form * (selecting code points from the Arabic Presentation Forms-B block). * @stable ICU 2.0 */ #define U_SHAPE_LETTERS_SHAPE_TASHKEEL_ISOLATED 0x18 /** Bit mask for letter shaping options. @stable ICU 2.0 */ #define U_SHAPE_LETTERS_MASK 0x18 /** Digit shaping option: do not perform digit shaping. @stable ICU 2.0 */ #define U_SHAPE_DIGITS_NOOP 0 /** * Digit shaping option: * Replace European digits (U+0030...) by Arabic-Indic digits. * @stable ICU 2.0 */ #define U_SHAPE_DIGITS_EN2AN 0x20 /** * Digit shaping option: * Replace Arabic-Indic digits by European digits (U+0030...). * @stable ICU 2.0 */ #define U_SHAPE_DIGITS_AN2EN 0x40 /** * Digit shaping option: * Replace European digits (U+0030...) by Arabic-Indic digits if the most recent * strongly directional character is an Arabic letter * (<code>u_charDirection()</code> result <code>U_RIGHT_TO_LEFT_ARABIC</code> [AL]).<br> * The direction of "preceding" depends on the direction indicator option. * For the first characters, the preceding strongly directional character * (initial state) is assumed to be not an Arabic letter * (it is <code>U_LEFT_TO_RIGHT</code> [L] or <code>U_RIGHT_TO_LEFT</code> [R]). * @stable ICU 2.0 */ #define U_SHAPE_DIGITS_ALEN2AN_INIT_LR 0x60 /** * Digit shaping option: * Replace European digits (U+0030...) by Arabic-Indic digits if the most recent * strongly directional character is an Arabic letter * (<code>u_charDirection()</code> result <code>U_RIGHT_TO_LEFT_ARABIC</code> [AL]).<br> * The direction of "preceding" depends on the direction indicator option. * For the first characters, the preceding strongly directional character * (initial state) is assumed to be an Arabic letter. * @stable ICU 2.0 */ #define U_SHAPE_DIGITS_ALEN2AN_INIT_AL 0x80 /** Not a valid option value. May be replaced by a new option. @stable ICU 2.0 */ #define U_SHAPE_DIGITS_RESERVED 0xa0 /** Bit mask for digit shaping options. @stable ICU 2.0 */ #define U_SHAPE_DIGITS_MASK 0xe0 /** Digit type option: Use Arabic-Indic digits (U+0660...U+0669). @stable ICU 2.0 */ #define U_SHAPE_DIGIT_TYPE_AN 0 /** Digit type option: Use Eastern (Extended) Arabic-Indic digits (U+06f0...U+06f9). @stable ICU 2.0 */ #define U_SHAPE_DIGIT_TYPE_AN_EXTENDED 0x100 /** Not a valid option value. May be replaced by a new option. @stable ICU 2.0 */ #define U_SHAPE_DIGIT_TYPE_RESERVED 0x200 /** Bit mask for digit type options. @stable ICU 2.0 */ #define U_SHAPE_DIGIT_TYPE_MASK 0x300 /* I need to change this from 0x3f00 to 0x300 */ /** * Tashkeel aggregation option: * Replaces any combination of U+0651 with one of * U+064C, U+064D, U+064E, U+064F, U+0650 with * U+FC5E, U+FC5F, U+FC60, U+FC61, U+FC62 consecutively. * @stable ICU 3.6 */ #define U_SHAPE_AGGREGATE_TASHKEEL 0x4000 /** Tashkeel aggregation option: do not aggregate tashkeels. @stable ICU 3.6 */ #define U_SHAPE_AGGREGATE_TASHKEEL_NOOP 0 /** Bit mask for tashkeel aggregation. @stable ICU 3.6 */ #define U_SHAPE_AGGREGATE_TASHKEEL_MASK 0x4000 /** * Presentation form option: * Don't replace Arabic Presentation Forms-A and Arabic Presentation Forms-B * characters with 0+06xx characters, before shaping. * @stable ICU 3.6 */ #define U_SHAPE_PRESERVE_PRESENTATION 0x8000 /** Presentation form option: * Replace Arabic Presentation Forms-A and Arabic Presentationo Forms-B with * their unshaped correspondents in range 0+06xx, before shaping. * @stable ICU 3.6 */ #define U_SHAPE_PRESERVE_PRESENTATION_NOOP 0 /** Bit mask for preserve presentation form. @stable ICU 3.6 */ #define U_SHAPE_PRESERVE_PRESENTATION_MASK 0x8000 /* Seen Tail option */ /** * Memory option: the result must have the same length as the source. * Shaping mode: The SEEN family character will expand into two characters using space near * the SEEN family character(i.e. the space after the character). * If there are no spaces found, an error U_NO_SPACE_AVAILABLE (as defined in utypes.h) * will be set in pErrorCode * * De-shaping mode: Any Seen character followed by Tail character will be * replaced by one cell Seen and a space will replace the Tail. * Affects: Seen options * @stable ICU 4.2 */ #define U_SHAPE_SEEN_TWOCELL_NEAR 0x200000 /** * Bit mask for Seen memory options. * @stable ICU 4.2 */ #define U_SHAPE_SEEN_MASK 0x700000 /* YehHamza option */ /** * Memory option: the result must have the same length as the source. * Shaping mode: The YEHHAMZA character will expand into two characters using space near it * (i.e. the space after the character * If there are no spaces found, an error U_NO_SPACE_AVAILABLE (as defined in utypes.h) * will be set in pErrorCode * * De-shaping mode: Any Yeh (final or isolated) character followed by Hamza character will be * replaced by one cell YehHamza and space will replace the Hamza. * Affects: YehHamza options * @stable ICU 4.2 */ #define U_SHAPE_YEHHAMZA_TWOCELL_NEAR 0x1000000 /** * Bit mask for YehHamza memory options. * @stable ICU 4.2 */ #define U_SHAPE_YEHHAMZA_MASK 0x3800000 /* New Tashkeel options */ /** * Memory option: the result must have the same length as the source. * Shaping mode: Tashkeel characters will be replaced by spaces. * Spaces will be placed at beginning of the buffer * * De-shaping mode: N/A * Affects: Tashkeel options * @stable ICU 4.2 */ #define U_SHAPE_TASHKEEL_BEGIN 0x40000 /** * Memory option: the result must have the same length as the source. * Shaping mode: Tashkeel characters will be replaced by spaces. * Spaces will be placed at end of the buffer * * De-shaping mode: N/A * Affects: Tashkeel options * @stable ICU 4.2 */ #define U_SHAPE_TASHKEEL_END 0x60000 /** * Memory option: allow the result to have a different length than the source. * Shaping mode: Tashkeel characters will be removed, buffer length will shrink. * De-shaping mode: N/A * * Affect: Tashkeel options * @stable ICU 4.2 */ #define U_SHAPE_TASHKEEL_RESIZE 0x80000 /** * Memory option: the result must have the same length as the source. * Shaping mode: Tashkeel characters will be replaced by Tatweel if it is connected to adjacent * characters (i.e. shaped on Tatweel) or replaced by space if it is not connected. * * De-shaping mode: N/A * Affects: YehHamza options * @stable ICU 4.2 */ #define U_SHAPE_TASHKEEL_REPLACE_BY_TATWEEL 0xC0000 /** * Bit mask for Tashkeel replacement with Space or Tatweel memory options. * @stable ICU 4.2 */ #define U_SHAPE_TASHKEEL_MASK 0xE0000 /* Space location Control options */ /** * This option affect the meaning of BEGIN and END options. if this option is not used the default * for BEGIN and END will be as following: * The Default (for both Visual LTR, Visual RTL and Logical Text) * 1. BEGIN always refers to the start address of physical memory. * 2. END always refers to the end address of physical memory. * * If this option is used it will swap the meaning of BEGIN and END only for Visual LTR text. * * The effect on BEGIN and END Memory Options will be as following: * A. BEGIN For Visual LTR text: This will be the beginning (right side) of the visual text( * corresponding to the physical memory address end for Visual LTR text, Same as END in * default behavior) * B. BEGIN For Logical text: Same as BEGIN in default behavior. * C. END For Visual LTR text: This will be the end (left side) of the visual text (corresponding * to the physical memory address beginning for Visual LTR text, Same as BEGIN in default behavior. * D. END For Logical text: Same as END in default behavior). * Affects: All LamAlef BEGIN, END and AUTO options. * @stable ICU 4.2 */ #define U_SHAPE_SPACES_RELATIVE_TO_TEXT_BEGIN_END 0x4000000 /** * Bit mask for swapping BEGIN and END for Visual LTR text * @stable ICU 4.2 */ #define U_SHAPE_SPACES_RELATIVE_TO_TEXT_MASK 0x4000000 /** * If this option is used, shaping will use the new Unicode code point for TAIL (i.e. 0xFE73). * If this option is not specified (Default), old unofficial Unicode TAIL code point is used (i.e. 0x200B) * De-shaping will not use this option as it will always search for both the new Unicode code point for the * TAIL (i.e. 0xFE73) or the old unofficial Unicode TAIL code point (i.e. 0x200B) and de-shape the * Seen-Family letter accordingly. * * Shaping Mode: Only shaping. * De-shaping Mode: N/A. * Affects: All Seen options * @stable ICU 4.8 */ #define U_SHAPE_TAIL_NEW_UNICODE 0x8000000 /** * Bit mask for new Unicode Tail option * @stable ICU 4.8 */ #define U_SHAPE_TAIL_TYPE_MASK 0x8000000 #endif