cannam@49: // Copyright (c) 2013-2014 Sandstorm Development Group, Inc. and contributors cannam@49: // Licensed under the MIT License: cannam@49: // cannam@49: // Permission is hereby granted, free of charge, to any person obtaining a copy cannam@49: // of this software and associated documentation files (the "Software"), to deal cannam@49: // in the Software without restriction, including without limitation the rights cannam@49: // to use, copy, modify, merge, publish, distribute, sublicense, and/or sell cannam@49: // copies of the Software, and to permit persons to whom the Software is cannam@49: // furnished to do so, subject to the following conditions: cannam@49: // cannam@49: // The above copyright notice and this permission notice shall be included in cannam@49: // all copies or substantial portions of the Software. cannam@49: // cannam@49: // THE SOFTWARE IS PROVIDED "AS IS", WITHOUT WARRANTY OF ANY KIND, EXPRESS OR cannam@49: // IMPLIED, INCLUDING BUT NOT LIMITED TO THE WARRANTIES OF MERCHANTABILITY, cannam@49: // FITNESS FOR A PARTICULAR PURPOSE AND NONINFRINGEMENT. IN NO EVENT SHALL THE cannam@49: // AUTHORS OR COPYRIGHT HOLDERS BE LIABLE FOR ANY CLAIM, DAMAGES OR OTHER cannam@49: // LIABILITY, WHETHER IN AN ACTION OF CONTRACT, TORT OR OTHERWISE, ARISING FROM, cannam@49: // OUT OF OR IN CONNECTION WITH THE SOFTWARE OR THE USE OR OTHER DEALINGS IN cannam@49: // THE SOFTWARE. cannam@49: cannam@49: #ifndef KJ_STRING_TREE_H_ cannam@49: #define KJ_STRING_TREE_H_ cannam@49: cannam@49: #if defined(__GNUC__) && !KJ_HEADER_WARNINGS cannam@49: #pragma GCC system_header cannam@49: #endif cannam@49: cannam@49: #include "string.h" cannam@49: cannam@49: namespace kj { cannam@49: cannam@49: class StringTree { cannam@49: // A long string, represented internally as a tree of strings. This data structure is like a cannam@49: // String, but optimized for concatenation and iteration at the expense of seek time. The cannam@49: // structure is intended to be used for building large text blobs from many small pieces, where cannam@49: // repeatedly concatenating smaller strings into larger ones would waste copies. This structure cannam@49: // is NOT intended for use cases requiring random access or computing substrings. For those, cannam@49: // you should use a Rope, which is a much more complicated data structure. cannam@49: // cannam@49: // The proper way to construct a StringTree is via kj::strTree(...), which works just like cannam@49: // kj::str(...) but returns a StringTree rather than a String. cannam@49: // cannam@49: // KJ_STRINGIFY() functions that construct large strings from many smaller strings are encouraged cannam@49: // to return StringTree rather than a flat char container. cannam@49: cannam@49: public: cannam@49: inline StringTree(): size_(0) {} cannam@49: inline StringTree(String&& text): size_(text.size()), text(kj::mv(text)) {} cannam@49: cannam@49: StringTree(Array&& pieces, StringPtr delim); cannam@49: // Build a StringTree by concatenating the given pieces, delimited by the given delimiter cannam@49: // (e.g. ", "). cannam@49: cannam@49: inline size_t size() const { return size_; } cannam@49: cannam@49: template cannam@49: void visit(Func&& func) const; cannam@49: cannam@49: String flatten() const; cannam@49: // Return the contents as a string. cannam@49: cannam@49: // TODO(someday): flatten() when *this is an rvalue and when branches.size() == 0 could simply cannam@49: // return `kj::mv(text)`. Requires reference qualifiers (Clang 3.3 / GCC 4.8). cannam@49: cannam@49: void flattenTo(char* __restrict__ target) const; cannam@49: // Copy the contents to the given character array. Does not add a NUL terminator. cannam@49: cannam@49: private: cannam@49: size_t size_; cannam@49: String text; cannam@49: cannam@49: struct Branch; cannam@49: Array branches; // In order. cannam@49: cannam@49: inline void fill(char* pos, size_t branchIndex); cannam@49: template cannam@49: void fill(char* pos, size_t branchIndex, First&& first, Rest&&... rest); cannam@49: template cannam@49: void fill(char* pos, size_t branchIndex, StringTree&& first, Rest&&... rest); cannam@49: template cannam@49: void fill(char* pos, size_t branchIndex, Array&& first, Rest&&... rest); cannam@49: template cannam@49: void fill(char* pos, size_t branchIndex, String&& first, Rest&&... rest); cannam@49: cannam@49: template cannam@49: static StringTree concat(Params&&... params); cannam@49: static StringTree&& concat(StringTree&& param) { return kj::mv(param); } cannam@49: cannam@49: template cannam@49: static inline size_t flatSize(const T& t) { return t.size(); } cannam@49: static inline size_t flatSize(String&& s) { return 0; } cannam@49: static inline size_t flatSize(StringTree&& s) { return 0; } cannam@49: cannam@49: template cannam@49: static inline size_t branchCount(const T& t) { return 0; } cannam@49: static inline size_t branchCount(String&& s) { return 1; } cannam@49: static inline size_t branchCount(StringTree&& s) { return 1; } cannam@49: cannam@49: template cannam@49: friend StringTree strTree(Params&&... params); cannam@49: }; cannam@49: cannam@49: inline StringTree&& KJ_STRINGIFY(StringTree&& tree) { return kj::mv(tree); } cannam@49: inline const StringTree& KJ_STRINGIFY(const StringTree& tree) { return tree; } cannam@49: cannam@49: inline StringTree KJ_STRINGIFY(Array&& trees) { return StringTree(kj::mv(trees), ""); } cannam@49: cannam@49: template cannam@49: StringTree strTree(Params&&... params); cannam@49: // Build a StringTree by stringifying the given parameters and concatenating the results. cannam@49: // If any of the parameters stringify to StringTree rvalues, they will be incorporated as cannam@49: // branches to avoid a copy. cannam@49: cannam@49: // ======================================================================================= cannam@49: // Inline implementation details cannam@49: cannam@49: namespace _ { // private cannam@49: cannam@49: template cannam@49: char* fill(char* __restrict__ target, const StringTree& first, Rest&&... rest) { cannam@49: // Make str() work with stringifiers that return StringTree by patching fill(). cannam@49: cannam@49: first.flattenTo(target); cannam@49: return fill(target + first.size(), kj::fwd(rest)...); cannam@49: } cannam@49: cannam@49: template constexpr bool isStringTree() { return false; } cannam@49: template <> constexpr bool isStringTree() { return true; } cannam@49: cannam@49: inline StringTree&& toStringTreeOrCharSequence(StringTree&& tree) { return kj::mv(tree); } cannam@49: inline StringTree toStringTreeOrCharSequence(String&& str) { return StringTree(kj::mv(str)); } cannam@49: cannam@49: template cannam@49: inline auto toStringTreeOrCharSequence(T&& value) cannam@49: -> decltype(toCharSequence(kj::fwd(value))) { cannam@49: static_assert(!isStringTree>(), cannam@49: "When passing a StringTree into kj::strTree(), either pass it by rvalue " cannam@49: "(use kj::mv(value)) or explicitly call value.flatten() to make a copy."); cannam@49: cannam@49: return toCharSequence(kj::fwd(value)); cannam@49: } cannam@49: cannam@49: } // namespace _ (private) cannam@49: cannam@49: struct StringTree::Branch { cannam@49: size_t index; cannam@49: // Index in `text` where this branch should be inserted. cannam@49: cannam@49: StringTree content; cannam@49: }; cannam@49: cannam@49: template cannam@49: void StringTree::visit(Func&& func) const { cannam@49: size_t pos = 0; cannam@49: for (auto& branch: branches) { cannam@49: if (branch.index > pos) { cannam@49: func(text.slice(pos, branch.index)); cannam@49: pos = branch.index; cannam@49: } cannam@49: branch.content.visit(func); cannam@49: } cannam@49: if (text.size() > pos) { cannam@49: func(text.slice(pos, text.size())); cannam@49: } cannam@49: } cannam@49: cannam@49: inline void StringTree::fill(char* pos, size_t branchIndex) { cannam@49: KJ_IREQUIRE(pos == text.end() && branchIndex == branches.size(), cannam@49: kj::str(text.end() - pos, ' ', branches.size() - branchIndex).cStr()); cannam@49: } cannam@49: cannam@49: template cannam@49: void StringTree::fill(char* pos, size_t branchIndex, First&& first, Rest&&... rest) { cannam@49: pos = _::fill(pos, kj::fwd(first)); cannam@49: fill(pos, branchIndex, kj::fwd(rest)...); cannam@49: } cannam@49: cannam@49: template cannam@49: void StringTree::fill(char* pos, size_t branchIndex, StringTree&& first, Rest&&... rest) { cannam@49: branches[branchIndex].index = pos - text.begin(); cannam@49: branches[branchIndex].content = kj::mv(first); cannam@49: fill(pos, branchIndex + 1, kj::fwd(rest)...); cannam@49: } cannam@49: cannam@49: template cannam@49: void StringTree::fill(char* pos, size_t branchIndex, String&& first, Rest&&... rest) { cannam@49: branches[branchIndex].index = pos - text.begin(); cannam@49: branches[branchIndex].content = StringTree(kj::mv(first)); cannam@49: fill(pos, branchIndex + 1, kj::fwd(rest)...); cannam@49: } cannam@49: cannam@49: template cannam@49: StringTree StringTree::concat(Params&&... params) { cannam@49: StringTree result; cannam@49: result.size_ = _::sum({params.size()...}); cannam@49: result.text = heapString( cannam@49: _::sum({StringTree::flatSize(kj::fwd(params))...})); cannam@49: result.branches = heapArray( cannam@49: _::sum({StringTree::branchCount(kj::fwd(params))...})); cannam@49: result.fill(result.text.begin(), 0, kj::fwd(params)...); cannam@49: return result; cannam@49: } cannam@49: cannam@49: template cannam@49: StringTree strTree(Params&&... params) { cannam@49: return StringTree::concat(_::toStringTreeOrCharSequence(kj::fwd(params))...); cannam@49: } cannam@49: cannam@49: } // namespace kj cannam@49: cannam@49: #endif // KJ_STRING_TREE_H_