|
|||||||||
| PREV CLASS NEXT CLASS | FRAMES NO FRAMES | ||||||||
| SUMMARY: NESTED | FIELD | CONSTR | METHOD | DETAIL: FIELD | CONSTR | METHOD | ||||||||
java.lang.Objectorg.basex.util.Token
public final class Token
This class provides convenience operations for handling so-called 'Tokens'. Tokens in BaseX are nothing else than UTF8 encoded strings, stored in a byte array. Note that, to guarantee a consistent string representation, all string conversions should be done via the methods of this class.
| Field Summary | |
|---|---|
static byte[] |
DOTS
Dots. |
static byte[] |
E_AMP
Ampersand Entity. |
static byte[] |
E_APOS
Apostrophe Entity. |
static byte[] |
E_CR
CarriageReturn Entity. |
static byte[] |
E_GT
GreaterThan Entity. |
static byte[] |
E_LT
LessThan Entity. |
static byte[] |
E_NL
NewLine Entity. |
static byte[] |
E_QU
Quote Entity. |
static byte[] |
E_TAB
Tab Entity. |
static byte[] |
EMPTY
Empty token. |
static byte[] |
FALSE
False token. |
static byte[] |
INF
Infinity. |
static int |
MAXLEN
Maximum length for hash calculation and index terms. |
static byte[] |
MININT
Minimum integer. |
static byte[] |
MZERO
Zero token. |
static byte[] |
NAN
Infinity. |
static byte[] |
NINF
Infinity. |
static byte[] |
ONE
One token. |
static byte[] |
SPACE
Space token. |
static byte[] |
TRUE
True token. |
static java.lang.String |
UTF8
UTF8 encoding string. |
static java.lang.String |
UTF82
UTF8 encoding string (variant). |
static byte[] |
ZERO
Zero token. |
| Method Summary | |
|---|---|
static byte[] |
append(byte[] t,
byte a)
Appends a character to the specified byte array. |
static byte[] |
append(byte[] t,
byte[] a)
Appends an array to the specified byte array. |
static boolean |
ascii(byte[] text)
Checks if the specified token only consists of ASCII characters. |
static byte[] |
chopNumber(byte[] t)
Finishes the numeric token, removing trailing zeroes. |
static int |
cl(int v)
Returns the expected codepoint length of the specified byte. |
static boolean |
contains(byte[] tok,
byte[] sub)
Checks if the first token contains the second token. |
static boolean |
contains(byte[] tok,
int c)
Checks if the first token contains the specified character. |
static boolean |
containslc(byte[] tok,
byte[] sub)
Checks if the first token contains the second token in lowercase. |
static int |
cp(byte[] b,
int i)
Returns the codepoint of the specified bytes, starting at the specified position. |
static byte[] |
delete(byte[] t,
byte[] c)
Deletes the specified characters out of the token. |
static byte[] |
delete(byte[] t,
int c)
Deletes the specified character out of the token. |
static int |
diff(byte[] tok,
byte[] tok2)
Calculates the difference of two character arrays. |
static boolean |
digit(int c)
Checks if the specified character is a digit. |
static boolean |
endsWith(byte[] tok,
byte[] sub)
Checks if the first token ends with the second token. |
static boolean |
endsWith(byte[] tok,
int c)
Checks if the first token starts with the specified character. |
static boolean |
eq(byte[] tok,
byte[] tok2)
Compares two character arrays for equality. |
static int |
ft(int ch)
Returns a lowercase ASCII character version of the specified character. |
static boolean |
ftChar(byte ch)
Checks if the specified character is a letter; special characters are converted to the standard ASCII charset. |
static boolean |
ftcontains(byte[] tok,
byte[] sub)
Checks if the first token contains the second fulltext term. |
static int |
ftNorm(int ch)
Returns a lowercase ASCII character version of the specified character. |
static int |
hash(byte[] tok)
Calculates a hash code for the specified token. |
static int |
indexOf(byte[] tok,
byte[] sub)
Returns the position of the specified token or -1. |
static int |
indexOf(byte[] tok,
byte[] sub,
int p)
Returns the position of the specified token or -1. |
static int |
indexOf(byte[] tok,
int c)
Returns the position of the specified character or -1. |
static byte[] |
insert(byte[] t,
int i,
byte[] c)
Inserts the specified characters at the specified position. |
static int |
lastIndexOf(byte[] tok,
int c)
Returns the last position of the specified character or -1. |
static byte[] |
lc(byte[] t)
Converts the specified token to lower case. |
static int |
lc(int ch)
Converts a character to lower case. |
static int |
len(byte[] text)
Returns the token length. |
static boolean |
letter(byte[] tok)
Checks if the specified token consists only of letters. |
static boolean |
letter(int c)
Checks if the specified character is a letter. |
static boolean |
letterOrDigit(byte[] tok)
Checks if the specified token consists only of letters and digits. |
static boolean |
letterOrDigit(int c)
Checks if the specified character is a letter or digit. |
static byte[] |
norm(byte[] tok)
Normalizes all whitespace occurrences from the specified token. |
static int |
numDigits(int x)
Checks number of digits of the specified integer. |
static byte[] |
replace(byte[] t,
int s,
int r)
Replaces the specified character and returns the result token. |
static byte[][] |
split(byte[] tok,
int sep)
Splits the token at all whitespaces and returns a array with all tokens. |
static boolean |
startsWith(byte[] tok,
byte[] sub)
Checks if the first token starts with the second token. |
static boolean |
startsWith(byte[] tok,
int c)
Checks if the first token starts with the specified character. |
static java.lang.String |
string(byte[] text)
Returns the specified token as string. |
static java.lang.String |
string(byte[] text,
int s,
int l)
Returns the specified token as string. |
static byte[] |
substring(byte[] tok,
int s)
Returns a subtoken of the specified token. |
static byte[] |
substring(byte[] tok,
int s,
int e)
Returns a subtoken of the specified token. |
static double |
toDouble(byte[] to)
Converts the specified token into a double value. |
static double |
toDouble(java.lang.String to)
Converts the specified string into a double value. |
static int |
toInt(byte[] to)
Converts the specified token into an integer value. |
static int |
toInt(byte[] to,
int ts,
int te)
Converts the specified token into an integer value. |
static int |
toInt(java.lang.String to)
Converts the specified string into an integer value. |
static byte[] |
token(boolean b)
Creates a byte array representation of the specified boolean value. |
static byte[] |
token(double d)
Creates a byte array representation from the specified double value; inspired by Xavier Franc's Qizx. |
static byte[] |
token(float f)
Creates a byte array representation from the specified float value. |
static byte[] |
token(int i)
Creates a byte array representation of the specified integer value. |
static byte[] |
token(long i)
Creates a byte array representation from the specified long value, using Java's standard method. |
static byte[] |
token(java.lang.String s)
Converts a string to a byte array. |
static long |
toLong(byte[] to)
Converts the specified token into an long value. |
static long |
toLong(byte[] to,
int ts,
int te)
Converts the specified token into an long value. |
static long |
toLong(java.lang.String to)
Converts the specified string into an long value. |
static byte[] |
translate(byte[] tok,
byte[] srch,
byte[] rep)
Performs a translation on the specified token. |
static byte[] |
trim(byte[] t)
Removes leading and trailing whitespaces from the specified token. |
static byte[] |
uc(byte[] t)
Converts the specified token to upper case. |
static int |
uc(int ch)
Switches a character to upper case. |
static byte[] |
utf8(byte[] s,
java.lang.String enc)
Converts a token from the input encoding to UTF8. |
static boolean |
ws(byte c)
Checks if the specified character is a whitespace. |
static boolean |
ws(byte[] tok)
Checks if the specified token has only whitespaces. |
| Methods inherited from class java.lang.Object |
|---|
equals, getClass, hashCode, notify, notifyAll, toString, wait, wait, wait |
| Field Detail |
|---|
public static final byte[] DOTS
public static final byte[] E_AMP
public static final byte[] E_APOS
public static final byte[] E_CR
public static final byte[] E_GT
public static final byte[] E_LT
public static final byte[] E_NL
public static final byte[] E_QU
public static final byte[] E_TAB
public static final byte[] EMPTY
public static final byte[] FALSE
public static final byte[] INF
public static final int MAXLEN
public static final byte[] MININT
public static final byte[] MZERO
public static final byte[] NAN
public static final byte[] NINF
public static final byte[] ONE
public static final byte[] SPACE
public static final byte[] TRUE
public static final java.lang.String UTF8
public static final java.lang.String UTF82
public static final byte[] ZERO
| Method Detail |
|---|
public static byte[] append(byte[] t,
byte a)
t - original tokena - token to be appended
public static byte[] append(byte[] t,
byte[] a)
t - original tokena - token to be appended
public static boolean ascii(byte[] text)
text - token
public static byte[] chopNumber(byte[] t)
t - token to be modified
public static int cl(int v)
v - first character byte
public static boolean contains(byte[] tok,
byte[] sub)
tok - first tokensub - second token
public static boolean contains(byte[] tok,
int c)
tok - first tokenc - character
public static boolean containslc(byte[] tok,
byte[] sub)
tok - first tokensub - second token
public static int cp(byte[] b,
int i)
b - byte arrayi - character position
public static byte[] delete(byte[] t,
byte[] c)
t - token to be checkedc - characters to be removed
public static byte[] delete(byte[] t,
int c)
t - token to be checkedc - character to be removed
public static int diff(byte[] tok,
byte[] tok2)
tok - token to be comparedtok2 - second token to be compared
public static boolean digit(int c)
c - the letter to be checked
public static boolean endsWith(byte[] tok,
byte[] sub)
tok - first tokensub - second token
public static boolean endsWith(byte[] tok,
int c)
tok - first tokenc - character
public static boolean eq(byte[] tok,
byte[] tok2)
tok - token to be comparedtok2 - second token to be compared
public static int ft(int ch)
ch - character to be converted
public static boolean ftChar(byte ch)
ch - character to be converted
public static boolean ftcontains(byte[] tok,
byte[] sub)
tok - first tokensub - second token
public static int ftNorm(int ch)
ch - character to be converted
public static int hash(byte[] tok)
tok - specified token
public static int indexOf(byte[] tok,
byte[] sub)
tok - first tokensub - second token
public static int indexOf(byte[] tok,
byte[] sub,
int p)
tok - first tokensub - second tokenp - start position
public static int indexOf(byte[] tok,
int c)
tok - first tokenc - character
public static byte[] insert(byte[] t,
int i,
byte[] c)
t - token to be modifiedi - insert positionc - characters to be inserted
public static int lastIndexOf(byte[] tok,
int c)
tok - first tokenc - character
public static byte[] lc(byte[] t)
t - token to be converted
public static int lc(int ch)
ch - character to be converted
public static int len(byte[] text)
text - token
public static boolean letter(byte[] tok)
tok - token to be checked
public static boolean letter(int c)
c - the letter to be checked
public static boolean letterOrDigit(byte[] tok)
tok - token to be checked
public static boolean letterOrDigit(int c)
c - the letter to be checked
public static byte[] norm(byte[] tok)
tok - token
public static int numDigits(int x)
x - number to be checked
public static byte[] replace(byte[] t,
int s,
int r)
t - token to be checkeds - the character to be replacedr - the new character
public static byte[][] split(byte[] tok,
int sep)
tok - token to be splitsep - separation character
public static boolean startsWith(byte[] tok,
byte[] sub)
tok - first tokensub - second token
public static boolean startsWith(byte[] tok,
int c)
tok - first tokenc - character
public static java.lang.String string(byte[] text)
text - token
public static java.lang.String string(byte[] text,
int s,
int l)
text - tokens - start positionl - length
public static byte[] substring(byte[] tok,
int s)
tok - tokens - start position
public static byte[] substring(byte[] tok,
int s,
int e)
tok - tokens - start positione - end position
public static double toDouble(byte[] to)
Double.NaN is returned when the input is invalid.
to - character array to be converted
public static double toDouble(java.lang.String to)
Double.NaN is returned when the input is invalid.
to - character array to be converted
public static int toInt(byte[] to)
Integer.MIN_VALUE is returned when the input is invalid.
to - character array to be converted
public static int toInt(byte[] to,
int ts,
int te)
Integer.MIN_VALUE is returned when the input is invalid.
to - character array to be convertedts - first byte to be parsedte - last byte to be parsed (exclusive)
public static int toInt(java.lang.String to)
Integer.MIN_VALUE is returned when the input is invalid.
to - character array to be converted
public static byte[] token(boolean b)
b - boolean value to be converted
public static byte[] token(double d)
d - double value to be converted
public static byte[] token(float f)
f - float value to be converted
public static byte[] token(int i)
i - int value to be converted
public static byte[] token(long i)
i - int value to be converted
public static byte[] token(java.lang.String s)
s - string to be converted
public static long toLong(byte[] to)
Long.MIN_VALUE is returned when the input is invalid.
to - character array to be converted
public static long toLong(byte[] to,
int ts,
int te)
Long.MIN_VALUE is returned when the input is invalid.
to - character array to be convertedts - first byte to be parsedte - last byte to be parsed - exclusive
public static long toLong(java.lang.String to)
Long.MIN_VALUE is returned when the input is invalid.
to - character array to be converted
public static byte[] translate(byte[] tok,
byte[] srch,
byte[] rep)
tok - tokensrch - characters to be foundrep - characters to be replaced
public static byte[] trim(byte[] t)
t - token to be checked
public static byte[] uc(byte[] t)
t - token to be converted
public static int uc(int ch)
ch - character to be converted
public static byte[] utf8(byte[] s,
java.lang.String enc)
s - token to be convertedenc - input encoding
public static boolean ws(byte c)
c - the letter to be checked
public static boolean ws(byte[] tok)
tok - token
|
|||||||||
| PREV CLASS NEXT CLASS | FRAMES NO FRAMES | ||||||||
| SUMMARY: NESTED | FIELD | CONSTR | METHOD | DETAIL: FIELD | CONSTR | METHOD | ||||||||