mirror of
				https://github.com/RGBCube/serenity
				synced 2025-10-25 07:32:32 +00:00 
			
		
		
		
	 5e1499d104
			
		
	
	
		5e1499d104
		
	
	
	
	
		
			
			This commit un-deprecates DeprecatedString, and repurposes it as a byte
string.
As the null state has already been removed, there are no other
particularly hairy blockers in repurposing this type as a byte string
(what it _really_ is).
This commit is auto-generated:
  $ xs=$(ack -l \bDeprecatedString\b\|deprecated_string AK Userland \
    Meta Ports Ladybird Tests Kernel)
  $ perl -pie 's/\bDeprecatedString\b/ByteString/g;
    s/deprecated_string/byte_string/g' $xs
  $ clang-format --style=file -i \
    $(git diff --name-only | grep \.cpp\|\.h)
  $ gn format $(git ls-files '*.gn' '*.gni')
		
	
			
		
			
				
	
	
		
			97 lines
		
	
	
	
		
			2.8 KiB
		
	
	
	
		
			C++
		
	
	
	
	
	
			
		
		
	
	
			97 lines
		
	
	
	
		
			2.8 KiB
		
	
	
	
		
			C++
		
	
	
	
	
	
| /*
 | |
|  * Copyright (c) 2020, Stephan Unverwerth <s.unverwerth@serenityos.org>
 | |
|  *
 | |
|  * SPDX-License-Identifier: BSD-2-Clause
 | |
|  */
 | |
| 
 | |
| #pragma once
 | |
| 
 | |
| #include "Token.h"
 | |
| 
 | |
| #include <AK/ByteString.h>
 | |
| #include <AK/HashMap.h>
 | |
| #include <AK/String.h>
 | |
| #include <AK/StringView.h>
 | |
| 
 | |
| namespace JS {
 | |
| 
 | |
| class Lexer {
 | |
| public:
 | |
|     explicit Lexer(StringView source, StringView filename = "(unknown)"sv, size_t line_number = 1, size_t line_column = 0);
 | |
| 
 | |
|     Token next();
 | |
| 
 | |
|     ByteString const& source() const { return m_source; }
 | |
|     String const& filename() const { return m_filename; }
 | |
| 
 | |
|     void disallow_html_comments() { m_allow_html_comments = false; }
 | |
| 
 | |
|     Token force_slash_as_regex();
 | |
| 
 | |
| private:
 | |
|     void consume();
 | |
|     bool consume_exponent();
 | |
|     bool consume_octal_number();
 | |
|     bool consume_hexadecimal_number();
 | |
|     bool consume_binary_number();
 | |
|     bool consume_decimal_number();
 | |
| 
 | |
|     bool is_unicode_character() const;
 | |
|     u32 current_code_point() const;
 | |
| 
 | |
|     bool is_eof() const;
 | |
|     bool is_line_terminator() const;
 | |
|     bool is_whitespace() const;
 | |
|     Optional<u32> is_identifier_unicode_escape(size_t& identifier_length) const;
 | |
|     Optional<u32> is_identifier_start(size_t& identifier_length) const;
 | |
|     Optional<u32> is_identifier_middle(size_t& identifier_length) const;
 | |
|     bool is_line_comment_start(bool line_has_token_yet) const;
 | |
|     bool is_block_comment_start() const;
 | |
|     bool is_block_comment_end() const;
 | |
|     bool is_numeric_literal_start() const;
 | |
|     bool match(char, char) const;
 | |
|     bool match(char, char, char) const;
 | |
|     bool match(char, char, char, char) const;
 | |
|     template<typename Callback>
 | |
|     bool match_numeric_literal_separator_followed_by(Callback) const;
 | |
|     bool slash_means_division() const;
 | |
| 
 | |
|     TokenType consume_regex_literal();
 | |
| 
 | |
|     ByteString m_source;
 | |
|     size_t m_position { 0 };
 | |
|     Token m_current_token;
 | |
|     char m_current_char { 0 };
 | |
|     bool m_eof { false };
 | |
| 
 | |
|     String m_filename;
 | |
|     size_t m_line_number { 1 };
 | |
|     size_t m_line_column { 0 };
 | |
| 
 | |
|     bool m_regex_is_in_character_class { false };
 | |
| 
 | |
|     struct TemplateState {
 | |
|         bool in_expr;
 | |
|         u8 open_bracket_count;
 | |
|     };
 | |
|     Vector<TemplateState> m_template_states;
 | |
| 
 | |
|     bool m_allow_html_comments { true };
 | |
| 
 | |
|     Optional<size_t> m_hit_invalid_unicode;
 | |
| 
 | |
|     static HashMap<DeprecatedFlyString, TokenType> s_keywords;
 | |
|     static HashMap<ByteString, TokenType> s_three_char_tokens;
 | |
|     static HashMap<ByteString, TokenType> s_two_char_tokens;
 | |
|     static HashMap<char, TokenType> s_single_char_tokens;
 | |
| 
 | |
|     struct ParsedIdentifiers : public RefCounted<ParsedIdentifiers> {
 | |
|         // Resolved identifiers must be kept alive for the duration of the parsing stage, otherwise
 | |
|         // the only references to these strings are deleted by the Token destructor.
 | |
|         HashTable<DeprecatedFlyString> identifiers;
 | |
|     };
 | |
| 
 | |
|     RefPtr<ParsedIdentifiers> m_parsed_identifiers;
 | |
| };
 | |
| 
 | |
| }
 |