Extensions/LatinWordExtensions.cs
(开头部分) 11KB这里只显示每个文件的开头 60 行。登录后可以解锁完整代码。
using System;
using System.Linq;
using System.Text;
namespace ComLib.Extensions
{
/// <summary>
/// Collection of extension utility methods useful for processing latin
/// text.
/// </summary>
public static class LatinWordExtensions
{
/// <summary>
/// A very liberal list of whitespace characters used in unicode.
/// Char.isWhitespace() only checks for latin1 values, this list is
/// more comprehensive and should not be used unintelligently for
/// stripping whitespace because it may make the target language
/// unreadable.
/// </summary>
public static readonly char[] WhiteSpaceCharacters = new char[]
{
'\x0020', // standard space
'\x0009', // horizontal tab
'\x000a', // LF
'\x000b', // line tabulation
'\x000c', // FF (form feed)
'\x000d', // CR
'\x00a0', // non-breaking space
'\x0085', // NEL (next line)
'\x1680', // OGHAM space mark
'\x2000', // En quad
'\x2001', // Em quad
'\x2002', // En space
'\x2003', // Em space
'\x2004', // three-per-em space
'\x2005', // four-per-em space
'\x2006', // six-per-em space
'\x2007', // figure space
'\x2008', // punctuation space
'\x2009', // thin space
'\x200a', // hair space
'\x200b', // zero width space
'\x2028', // LS (unicode line separator)
'\x2029', // paragraph separator
'\x202f', // narrow no-break space
'\x205f', // medium mathematical space
'\x3000', // ideographic space (CJKV)
};
/// <summary>
/// Small list of english words that commonly excluded from abbreviations
/// and initials.
/// </summary>
public static readonly string[] EnglishInitialsExclusionWords = new string[]
{
"a", "an", "as", "at", "by", "for", "from", "in", "into", "of", "on",
"to", "the", "with"
};
/// <summary>
后面还有 245 行代码,解锁后查看完整代码
24 小时内免费解锁 3 个项目,之后 1 积分/个。 规则说明
AI 解读
登录后可用,每次 10 积分,解读结果公开显示在下面。
还没有人解读过这个文件。
